{"id":"W2118786685","doi":"10.1111/j.1365-2929.2006.02541.x","title":"Assessment of clinical reasoning in the context of uncertainty: the effect of variability within the reference panel","year":2006,"lang":"en","type":"article","venue":"Medical Education","topic":"Clinical Reasoning and Diagnostic Skills","field":"Medicine","cited_by":68,"is_retracted":false,"has_abstract":true,"ca_institutions":"Université de Montréal","funders":"","keywords":"Concordance; Context (archaeology); Cronbach's alpha; Test (biology); Statistics; Psychology; Medicine; Clinical psychology; Mathematics; Psychometrics; Internal medicine","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.01890254,0.0000922687,0.0004768094,0.00002284858,0.00003313888,0.000004394352,0.0002826641,0.0001529524,0.00009031202],"category_scores_gemma":[0.2596491,0.00003636197,0.0001236659,0.0002591247,0.0006727099,0.00001407962,0.00003471679,0.0005949952,0.000001139981],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.00003193709,"about_ca_system_score_gemma":0.001704188,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.004256749,"about_ca_topic_score_gemma":0.0001838802,"domain_scores_codex":[0.9958044,0.001894502,0.001179637,0.0001922524,0.0008093898,0.0001197808],"domain_scores_gemma":[0.8672177,0.1303018,0.0009523037,0.001091919,0.0003302941,0.0001060546],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"observational","study_design_scores_codex":[0.0001772281,0.001315831,0.924513,0.00009323173,0.0000281645,0.000001098922,0.0005433111,0.00002878718,0.00001935803,0.01271307,0.002383978,0.05818295],"study_design_scores_gemma":[0.001095826,0.0006249322,0.9917542,0.001525737,0.000140002,0.000007637822,0.0009509469,0.002703846,0.00009079675,0.00078304,0.0002851585,0.00003791409],"study_design_candidate":"observational","study_design_consensus":"observational","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9861162,0.000114429,0.000117452,0.009241363,0.0003559331,0.0005031687,0.000003351676,0.000005210047,0.003542933],"genre_scores_gemma":[0.9985113,0.00002826704,0.0001310402,0.001022022,0.0001807942,0.00005536639,0.0000321089,0.000004740671,0.00003440837],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.2407466,"threshold_uncertainty_score":0.7465872,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.03606216007835677,"score_gpt":0.436004507679344,"score_spread":0.3999423476009872,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}