{"id":"W2188660453","doi":"10.4300/jgme-d-15-00437.1","title":"General Versus Technique-Specific Surgical Skills Assessments: Do We Need to Reinvent the Wheel?","year":2015,"lang":"en","type":"article","venue":"Journal of Graduate Medical Education","topic":"Innovations in Medical Education","field":"Medicine","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Manitoba","funders":"","keywords":"Cronbach's alpha; Construct validity; Reliability (semiconductor); Rating scale; Concurrent validity; Medicine; Laparoscopic cholecystectomy; Construct (python library); Educational measurement; Scale (ratio); Medical physics; Psychometrics; Medical education; Psychology; Curriculum; Surgery; Clinical psychology; Computer science; Internal consistency","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00421074,0.0002022428,0.0004027777,0.0004647782,0.0001110459,0.00006428905,0.0004277987,0.0002118563,0.0004488493],"category_scores_gemma":[0.005136749,0.0001276425,0.0001535482,0.001208978,0.0001975777,0.000186139,0.0000645753,0.001183482,0.00007987807],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0008948651,"about_ca_system_score_gemma":0.004831117,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00001833964,"about_ca_topic_score_gemma":0.000001069507,"domain_scores_codex":[0.9952276,0.0002519425,0.001138821,0.0002288873,0.002830818,0.0003219062],"domain_scores_gemma":[0.996271,0.0001545982,0.0004911544,0.0004572767,0.001890016,0.0007359935],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","study_design_scores_codex":[0.0005606065,0.001623406,0.0003384912,0.00003096788,0.00007864985,0.00004899217,0.001023634,0.00000767136,0.0001270685,0.003271634,0.7259288,0.2669601],"study_design_scores_gemma":[0.003572774,0.001468954,0.002170477,0.0009130504,0.0001428336,0.001631069,0.004036984,0.0001685147,0.0004815096,0.002057648,0.9831499,0.0002062175],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.5401155,0.00136056,0.005736009,0.4340317,0.01473781,0.001012889,0.000001498016,0.00004646201,0.002957602],"genre_scores_gemma":[0.9043369,0.002343385,0.06161565,0.01225572,0.01591742,0.0002952152,0.0001003566,0.0001041217,0.003031228],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.421776,"threshold_uncertainty_score":0.8570193,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.07396001131527487,"score_gpt":0.4201404088972802,"score_spread":0.3461803975820054,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}