{"id":"W2063009847","doi":"10.1002/ir.4","title":"Improving Judgments About Teaching Effectiveness Using Teacher Rating Forms","year":2001,"lang":"en","type":"article","venue":"New Directions for Institutional Research","topic":"Evaluation of Teaching Practices","field":"Social Sciences","cited_by":85,"is_retracted":false,"has_abstract":true,"ca_institutions":"Concordia University","funders":"","keywords":"Summative assessment; Process (computing); Computer science; Evaluation methods; Psychology; Management science; Mathematics education; Process management; Formative assessment; Business; Engineering","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.1337067,0.0009321541,0.00215387,0.005262025,0.0009848151,0.005698043,0.001900471,0.0009752659,0.003423055],"category_scores_gemma":[0.4348013,0.000416196,0.001012209,0.002475017,0.001363647,0.004717635,0.002132553,0.002333746,0.001493165],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001604725,"about_ca_system_score_gemma":0.002046502,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001644365,"about_ca_topic_score_gemma":0.00312872,"domain_scores_codex":[0.7919956,0.152095,0.01533963,0.003049772,0.03598813,0.001531846],"domain_scores_gemma":[0.397953,0.450886,0.034859,0.01928448,0.09419633,0.002821244],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"observational","study_design_scores_codex":[0.001219561,0.0008854871,0.1157515,0.001168227,0.0004067878,0.00008044798,0.01076461,0.003917505,0.005565427,0.005695011,0.01003079,0.8445145],"study_design_scores_gemma":[0.0009014704,0.01236312,0.6238843,0.006000064,0.001559288,0.0008273036,0.04038846,0.1106624,0.09700325,0.03793896,0.06691449,0.001556953],"study_design_candidate":"observational","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.7410725,0.00243346,0.1911688,0.004285146,0.0007318821,0.002480035,0.0009126228,0.002098175,0.05481734],"genre_scores_gemma":[0.8852738,0.001042138,0.1093162,0.0003034634,0.0002223049,0.0009156878,0.0003891047,0.0001532252,0.002384189],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.1337067,"threshold_uncertainty_score":0.707117,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.362434530184532,"score_gpt":0.5635661835029041,"score_spread":0.2011316533183721,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}