{"id":"W2171297463","doi":"10.1016/s0140-6736(14)60376-7","title":"Bias in assessing trainees' clinical competence: the influence of assessors' recent experiences of other performances on present assessment scores","year":2014,"lang":"en","type":"article","venue":"The Lancet","topic":"Clinical Reasoning and Diagnostic Skills","field":"Medicine","cited_by":0,"is_retracted":false,"has_abstract":false,"ca_institutions":"University of British Columbia; Dalhousie University","funders":"","keywords":"Competence (human resources); Confidence interval; Psychology; Medicine; Clinical psychology; Social psychology","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.272159,0.0005572734,0.001474862,0.004329853,0.001646558,0.003124629,0.001929826,0.002268008,0.001197902],"category_scores_gemma":[0.5928326,0.0005765279,0.002464794,0.00343933,0.003247638,0.003790481,0.003753632,0.001822758,0.0003142408],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001570236,"about_ca_system_score_gemma":0.002067983,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.005247459,"about_ca_topic_score_gemma":0.0103788,"domain_scores_codex":[0.7235535,0.191965,0.03691529,0.01162592,0.03266391,0.00327629],"domain_scores_gemma":[0.1825018,0.7032555,0.05769351,0.0325569,0.02048806,0.003504082],"domain_codex":null,"domain_gemma":"evaluation","domain_candidate":"evaluation","domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"observational","study_design_scores_codex":[0.001645633,0.00008461166,0.957221,0.0003176578,0.001836079,0.00009440679,0.00594255,0.0004611686,0.0005145019,0.0004636184,0.0005855116,0.03083319],"study_design_scores_gemma":[0.0001556613,0.001284033,0.9836632,0.0005450681,0.001056654,0.001057331,0.002159654,0.003728202,0.00164238,0.002237962,0.002320748,0.0001491281],"study_design_candidate":"observational","study_design_consensus":"observational","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9700691,0.007461274,0.01309085,0.002165332,0.0005258425,0.0001648782,0.0002791791,0.00007945139,0.006164214],"genre_scores_gemma":[0.9963487,0.0002315347,0.002460219,0.0002981102,0.0001698182,0.00006804058,0.00009661297,0.0000414535,0.0002853733],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.727841,"threshold_uncertainty_score":0.8975576,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.1333152110856842,"score_gpt":0.4544746419643681,"score_spread":0.3211594308786839,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}