{"id":"W4304688367","doi":"10.1007/s40670-022-01654-2","title":"Correlation of Narrative Evaluations to Clerkship Grades Using Statistical Sentiment Analysis","year":2022,"lang":"en","type":"article","venue":"Medical Science Educator","topic":"Innovations in Medical Education","field":"Medicine","cited_by":5,"is_retracted":false,"has_abstract":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Narrative; Sentiment analysis; Character (mathematics); Psychology; Word (group theory); Mathematics education; Computer science; Linguistics; Natural language processing; Mathematics","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.003822498,0.00009809293,0.0002671918,0.0009642304,0.0006464816,0.00001820504,0.0002899498,0.0000426882,0.01167877],"category_scores_gemma":[0.007564094,0.00008855757,0.0000627054,0.008461454,0.0007309784,0.00009631928,0.0001453907,0.0003587219,0.00001683069],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0008754963,"about_ca_system_score_gemma":0.004355311,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.000125056,"about_ca_topic_score_gemma":0.000005510175,"domain_scores_codex":[0.9951122,0.0001481311,0.0005279377,0.0003868977,0.003514143,0.0003106895],"domain_scores_gemma":[0.998492,0.0001594169,0.0001490469,0.0003196468,0.0004616014,0.0004182807],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"observational","study_design_scores_codex":[0.0003447394,0.01009059,0.5447187,0.0001907149,0.001069357,0.0000351847,0.1242558,0.01426436,0.01869873,0.06574669,0.06513312,0.155452],"study_design_scores_gemma":[0.001134134,0.0009025022,0.5432806,0.0001089145,0.001529396,0.00008607768,0.06654149,0.3784387,0.001500494,0.001425196,0.00462379,0.0004287542],"study_design_candidate":"observational","study_design_consensus":"observational","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.94187,0.00001606875,0.04440035,0.01170959,0.001310306,0.0004201821,0.000009312267,0.00001968821,0.000244476],"genre_scores_gemma":[0.9820743,5.039662e-7,0.01587885,0.001539108,0.0001591783,0.00009541976,0.0000571174,0.000007166134,0.000188368],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.3641743,"threshold_uncertainty_score":0.9892247,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.05160002977890953,"score_gpt":0.4587432759342102,"score_spread":0.4071432461553006,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}