{"id":"W4385852392","doi":"10.53300/001c.86151","title":"Student Evaluations of Teaching: Understanding Limitations and Advocating for a Gold Standard for Measuring Teaching Effectiveness","year":2023,"lang":"en","type":"article","venue":"Legal Education Review","topic":"Legal Issues in Education","field":"Social Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Promotion (chess); Set (abstract data type); Likert scale; Psychology; Mathematics education; Value (mathematics); Scale (ratio); Medical education; Public relations; Sociology; Computer science; Political science; Medicine; Law","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.5796067,0.001610254,0.00365409,0.01204034,0.006338612,0.01935677,0.01059367,0.01536385,0.00244531],"category_scores_gemma":[0.651567,0.001136319,0.002431337,0.008341606,0.03176695,0.01758463,0.01964209,0.02672847,0.001927382],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.009418602,"about_ca_system_score_gemma":0.01470545,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.007096254,"about_ca_topic_score_gemma":0.007437453,"domain_scores_codex":[0.4657891,0.3180721,0.05618652,0.02328522,0.1317821,0.004885004],"domain_scores_gemma":[0.2792627,0.503239,0.03418468,0.05365563,0.1236618,0.005996226],"domain_codex":null,"domain_gemma":"evaluation","domain_candidate":"evaluation","domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.0003548074,0.0002078392,0.05576054,0.003037332,0.0004828481,0.0002183167,0.02551899,0.0006708898,0.0008783549,0.387408,0.2076503,0.3178118],"study_design_scores_gemma":[0.0002254057,0.0008716556,0.04532275,0.0235946,0.0003153115,0.0008564882,0.01951768,0.01001386,0.003793859,0.4618972,0.4329425,0.0006486865],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"commentary","genre_gemma":"empirical","genre_scores_codex":[0.0220124,0.02080392,0.1378096,0.7578583,0.02519056,0.001121638,0.0008248655,0.0005462923,0.03383239],"genre_scores_gemma":[0.4408602,0.006957399,0.2647618,0.260078,0.01275246,0.00548643,0.0007438135,0.0009228147,0.007437099],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.4203933,"threshold_uncertainty_score":0.5184199,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.2123780634492837,"score_gpt":0.5089079936489888,"score_spread":0.2965299301997051,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}