{"id":"W2083506651","doi":"10.1186/1472-6920-8-58","title":"A generalizability study of the medical judgment vignettes interview to assess students' noncognitive attributes for medical school","year":2008,"lang":"en","type":"article","venue":"BMC Medical Education","topic":"Medical Education and Admissions","field":"Medicine","cited_by":22,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Calgary","funders":"","keywords":"Generalizability theory; Rubric; Vignette; Reliability (semiconductor); Psychology; Interview; Applied psychology; Medical education; Social psychology; Clinical psychology; Medicine; Developmental psychology; Mathematics education","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaresearch","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.003685353,0.0002665911,0.0006783006,0.0001424204,0.0002451649,0.00001633648,0.0009090058,0.0003517769,0.02708193],"category_scores_gemma":[0.2199005,0.0001661893,0.0002444739,0.000651467,0.0004523091,0.00006101266,0.0003444479,0.0006337877,0.00005275187],"about_ca_system_candidate":true,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0002322267,"about_ca_system_score_gemma":0.03666305,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.000254845,"about_ca_topic_score_gemma":0.0002245362,"domain_scores_codex":[0.9915551,0.000861296,0.001263123,0.0005911778,0.005329606,0.0003996892],"domain_scores_gemma":[0.9860656,0.002190041,0.0003167522,0.0009647661,0.001063185,0.009399665],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"observational","study_design_scores_codex":[0.0003556669,0.0523069,0.4937054,0.00131155,0.0003218227,0.00002285933,0.005013662,0.000001128833,0.00004842933,0.0004646653,0.3995584,0.04688951],"study_design_scores_gemma":[0.01367106,0.003165534,0.8949412,0.007186346,0.0006662106,0.0004419947,0.01748791,0.001072028,0.0008076783,0.00031868,0.059535,0.0007063402],"study_design_candidate":"observational","study_design_consensus":"observational","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.8555114,0.0002329308,0.00627909,0.1320868,0.00217933,0.003400437,0.00001037616,0.00005718833,0.0002425305],"genre_scores_gemma":[0.9271599,0.0001577463,0.001483815,0.06666926,0.001381892,0.001508005,0.00009221334,0.00003497799,0.001512186],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.4012358,"threshold_uncertainty_score":0.9738075,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.1500344980151584,"score_gpt":0.4608750537801394,"score_spread":0.3108405557649809,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}