{"id":"W2207150368","doi":"10.1016/j.jsurg.2015.04.020","title":"Reliable Assessment of Performance in Surgery: A Practical Approach to Generalizability Theory","year":2015,"lang":"en","type":"editorial","venue":"Journal of surgical education","topic":"Surgical Simulation and Training","field":"Medicine","cited_by":8,"is_retracted":false,"has_abstract":false,"ca_institutions":"McGill University","funders":"","keywords":"Generalizability theory; Debriefing; Preparedness; Thematic analysis; Medical education; Perioperative; Medicine; Psychology; Qualitative research; Surgery","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.05851856,0.003527135,0.005613983,0.005449769,0.00384088,0.01032016,0.006283789,0.01993771,0.004255557],"category_scores_gemma":[0.2762808,0.001687706,0.003953291,0.002290593,0.00991645,0.008068839,0.002577066,0.03010138,0.002481416],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.003122823,"about_ca_system_score_gemma":0.006714465,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.002216869,"about_ca_topic_score_gemma":0.006010708,"domain_scores_codex":[0.958405,0.01533084,0.00859811,0.003277699,0.01384931,0.000539131],"domain_scores_gemma":[0.5325139,0.3947875,0.00611859,0.008192362,0.05532319,0.003064523],"domain_codex":null,"domain_gemma":"evaluation","domain_candidate":"evaluation","domain_consensus":null,"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","study_design_scores_codex":[0.0000551304,0.00003107904,0.0002369087,0.001682335,0.000213184,0.0002651059,0.0001410923,0.00009520839,0.0000818983,0.003590865,0.9667481,0.02685916],"study_design_scores_gemma":[0.0001887152,0.0001069611,0.003055996,0.004950549,0.0007831038,0.001084956,0.0006436112,0.001490827,0.0004541018,0.02509549,0.961956,0.0001896234],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","genre_codex":"editorial","genre_gemma":"editorial","genre_scores_codex":[0.0000930821,0.01088582,0.002000999,0.05401223,0.9320808,0.00004034799,0.00007215019,0.00004297373,0.0007715619],"genre_scores_gemma":[0.001908287,0.007188243,0.002404117,0.02160933,0.9642102,0.0001010387,0.00003549563,0.00006183939,0.002481537],"genre_candidate":"editorial","genre_consensus":"editorial","teacher_disagreement_score":0.9414815,"threshold_uncertainty_score":0.3094794,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.05541464794960498,"score_gpt":0.4063985187289122,"score_spread":0.3509838707793072,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}