{"id":"W2471842352","doi":"10.2310/8000.2011.110398","title":"How do I improve the quality of in-training assessment of learners?","year":2011,"lang":"en","type":"article","venue":"Canadian Journal of Emergency Medicine","topic":"Innovations in Medical Education","field":"Medicine","cited_by":4,"is_retracted":false,"has_abstract":false,"ca_institutions":"St. Michael's Hospital","funders":"","keywords":"CLARITY; Medicine; Consistency (knowledge bases); Medical education; Quality (philosophy); Honesty; Realm; Process (computing); Psychology; Computer science; Artificial intelligence; Social psychology","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02150476,0.0004891043,0.00115958,0.0009677482,0.001236244,0.004741042,0.001699492,0.002564729,0.005186868],"category_scores_gemma":[0.1987505,0.0003330298,0.001074947,0.0008740407,0.001190798,0.004976703,0.002617361,0.002569783,0.001934184],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001919676,"about_ca_system_score_gemma":0.006317214,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.004155133,"about_ca_topic_score_gemma":0.00750706,"domain_scores_codex":[0.9861264,0.00788131,0.0009365326,0.001257671,0.002857618,0.0009404173],"domain_scores_gemma":[0.892952,0.05486605,0.01504899,0.007527771,0.0191716,0.01043353],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"not_applicable","study_design_scores_codex":[0.00143091,0.004709967,0.1831589,0.001512193,0.0005253005,0.0001917786,0.008701245,0.0007271632,0.001514877,0.001913487,0.069237,0.7263771],"study_design_scores_gemma":[0.002732748,0.01296858,0.6643937,0.01615733,0.003485769,0.002940033,0.04445504,0.006677972,0.02348268,0.03208244,0.1896767,0.0009470737],"study_design_candidate":"not_applicable","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.6261504,0.01565535,0.02114789,0.2547569,0.007907254,0.0007245468,0.000976624,0.001341496,0.07133956],"genre_scores_gemma":[0.9302623,0.008963871,0.03537431,0.01698765,0.00210441,0.0005919574,0.0003814558,0.0002788343,0.005055351],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.02150476,"threshold_uncertainty_score":0.1137294,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.2294963020615004,"score_gpt":0.4432331484250108,"score_spread":0.2137368463635104,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}