{"id":"W3008964715","doi":"10.1007/s12630-020-01597-5","title":"The lack of construct validity when assessing clinical clerks during their anesthesia rotations","year":2020,"lang":"en","type":"article","venue":"Canadian Journal of Anesthesia/Journal canadien d anesthésie","topic":"Cardiac, Anesthesia and Surgical Outcomes","field":"Medicine","cited_by":1,"is_retracted":false,"has_abstract":false,"ca_institutions":"McMaster University","funders":"","keywords":"Construct validity; Construct (python library); Psychology; Medicine; Anesthesia; Clinical psychology; Computer science; Psychometrics","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.07806344,0.0007075773,0.00112246,0.004026241,0.003365958,0.005025013,0.001977081,0.001334357,0.002062379],"category_scores_gemma":[0.2093947,0.0008795178,0.002468723,0.002857641,0.003864676,0.002567674,0.004484004,0.002646127,0.0006555779],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.005401835,"about_ca_system_score_gemma":0.01082836,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.01603685,"about_ca_topic_score_gemma":0.04350988,"domain_scores_codex":[0.9170277,0.03151222,0.01300005,0.004409619,0.03028567,0.00376483],"domain_scores_gemma":[0.6863847,0.2017417,0.03142122,0.01684045,0.05575364,0.007858216],"domain_codex":null,"domain_gemma":"methods","domain_candidate":"methods","domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"observational","study_design_scores_codex":[0.000183738,0.0002868358,0.9544482,0.0003135857,0.0004641096,0.00006807121,0.009859563,0.0003429493,0.0004938178,0.001544919,0.002002167,0.02999213],"study_design_scores_gemma":[0.00005850306,0.0004726889,0.9658911,0.0008930948,0.0001689691,0.0002161481,0.0200845,0.002529426,0.0009425601,0.001910777,0.0067391,0.00009300352],"study_design_candidate":"observational","study_design_consensus":"observational","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9713756,0.00134895,0.00527579,0.002213074,0.0007700762,0.0005862063,0.0004652896,0.0000597059,0.01790529],"genre_scores_gemma":[0.993091,0.0003519298,0.003911329,0.0006096282,0.00009852178,0.0004645855,0.0004523016,0.00004195181,0.0009787807],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.9219366,"threshold_uncertainty_score":0.4128439,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.07170405453465867,"score_gpt":0.3071154899280361,"score_spread":0.2354114353933774,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}