{"id":"W2163707981","doi":"10.1016/j.gie.2015.03.1450","title":"Su1553 Can Novice Endoscopists Accurately Self-Assess Performance During Their Initial Clinical Colonoscopies? a Prospective, Cross-Sectional Study","year":2015,"lang":"en","type":"article","venue":"Gastrointestinal Endoscopy","topic":"Reflective Practices in Education","field":"Social Sciences","cited_by":2,"is_retracted":false,"has_abstract":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Medicine; Colonoscopy; Endoscopy; Self-assessment; Medical physics; Cross-sectional study; Physical therapy; Radiology; Pathology; Internal medicine; Colorectal cancer; Social psychology","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaepi_narrow","sts"],"consensus_categories":[],"category_scores_codex":[0.004443751,0.000424046,0.0004736367,0.0002048485,0.001300429,0.0006918876,0.0008404654,0.0001306778,0.00008030437],"category_scores_gemma":[0.006647526,0.00042571,0.0001267643,0.0008186233,0.0006886621,0.001519713,0.0002417873,0.001044059,0.0001346285],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.00169253,"about_ca_system_score_gemma":0.003517333,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.002630307,"about_ca_topic_score_gemma":0.002514203,"domain_scores_codex":[0.9941121,0.001790018,0.0009993656,0.0009836109,0.001125949,0.0009889377],"domain_scores_gemma":[0.9951783,0.001496365,0.0007662348,0.0004769886,0.001466974,0.0006151027],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"observational","study_design_scores_codex":[0.001129302,0.001469388,0.987667,0.00002889342,0.0000593587,0.00002904567,0.008709758,0.00005891,0.0004446658,0.0002538247,0.00005007079,0.00009974407],"study_design_scores_gemma":[0.002871085,0.002027244,0.9859353,0.00007189944,0.00003423323,0.0003709838,0.007131653,0.0000763394,0.0005637773,0.0002685761,0.0003157295,0.0003332049],"study_design_candidate":"observational","study_design_consensus":"observational","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9715996,0.00002171675,0.0002214407,0.0001975279,0.002308804,0.001590273,0.00002061403,0.0004366294,0.02360343],"genre_scores_gemma":[0.9783533,0.00001587122,0.01891911,0.00007175314,0.001979925,0.0003401663,0.0000111664,0.00005300765,0.0002557004],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.02334773,"threshold_uncertainty_score":0.9999998,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.144413679828016,"score_gpt":0.4723894625128158,"score_spread":0.3279757826847998,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}