{"id":"W3137392471","doi":"10.1007/s12630-021-01971-x","title":"Investigating faculty assessment of anesthesia trainees and the failing-to-fail phenomenon: a randomized controlled trial","year":2021,"lang":"en","type":"article","venue":"Canadian Journal of Anesthesia/Journal canadien d anesthésie","topic":"Innovations in Medical Education","field":"Medicine","cited_by":7,"is_retracted":false,"has_abstract":false,"ca_institutions":"Children's Hospital of Eastern Ontario; Hospital for Sick Children; University of Ottawa; Sunnybrook Health Science Centre; Sinai Health System; SickKids Foundation; University of Toronto","funders":"Canadian Anesthesiologists' Society","keywords":"Affect (linguistics); Randomized controlled trial; Medicine; Psychology; Medical education; Internal medicine","routes":{"ca_aff":true,"ca_fund":true,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.009258023,0.002465995,0.00581462,0.0007280192,0.001754518,0.002183031,0.002049146,0.007738609,0.01125542],"category_scores_gemma":[0.02299522,0.001719309,0.00454857,0.001028142,0.003142696,0.002870162,0.001221829,0.005738266,0.00133039],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001623381,"about_ca_system_score_gemma":0.002957232,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.002514691,"about_ca_topic_score_gemma":0.003390484,"domain_scores_codex":[0.9896803,0.006461091,0.0009570876,0.001718804,0.000538116,0.0006445855],"domain_scores_gemma":[0.9746487,0.01419606,0.005852859,0.001545047,0.00120842,0.002548954],"domain_codex":null,"domain_gemma":"evaluation","domain_candidate":"evaluation","domain_consensus":null,"study_design_codex":"randomized_trial","study_design_gemma":"randomized_trial","study_design_scores_codex":[0.9735991,0.01035986,0.004803382,0.001528719,0.004178911,0.00004521067,0.0002041157,0.0001299253,0.0002898345,0.0001362993,0.0004102653,0.004314347],"study_design_scores_gemma":[0.8484513,0.1334006,0.009492642,0.0004276637,0.006123281,0.00003831782,0.0003235447,0.0005788858,0.0001839927,0.0003943569,0.0005172013,0.00006821743],"study_design_candidate":"randomized_trial","study_design_consensus":"randomized_trial","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9781359,0.008343347,0.0009985932,0.0008746639,0.001579907,0.007898098,0.000509574,0.00005790202,0.00160208],"genre_scores_gemma":[0.9834502,0.001952629,0.001926031,0.001223261,0.0008058696,0.009168931,0.0003171033,0.00001276488,0.001143266],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.990742,"threshold_uncertainty_score":0.04896164,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.01734286607969217,"score_gpt":0.2874794096117828,"score_spread":0.2701365435320907,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}