{"id":"W2515878609","doi":"","title":"Book Review: A. L. LAVIGNE & T. L. GOOD. (2014). Teacher and Student Evaluation: Moving Beyond the Failure of School Reform.","year":2015,"lang":"en","type":"article","venue":"McGill Journal of Education / Revue des sciences de l'éducation de McGill","topic":"Evaluation and Performance Assessment","field":"Decision Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"McGill University","funders":"","keywords":"Political science; Sociology; Management; Economics","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaresearch","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.03436421,0.0002018243,0.0003989303,0.0004911356,0.0008850477,0.0001978136,0.001201932,0.00008425485,0.002254922],"category_scores_gemma":[0.005558663,0.0001301319,0.0001436483,0.001249379,0.0004936895,0.001821662,0.0001167665,0.0003189516,0.00006547994],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.002031695,"about_ca_system_score_gemma":0.003144452,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00005835354,"about_ca_topic_score_gemma":0.00009032527,"domain_scores_codex":[0.9937539,0.001290637,0.001580059,0.0003726881,0.002689119,0.0003135681],"domain_scores_gemma":[0.9919313,0.0004659586,0.002042175,0.0005641034,0.004575725,0.0004207318],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","study_design_scores_codex":[0.00004261394,0.0007525439,0.0197106,0.000192629,0.00007171196,0.000001322608,0.006682243,0.003469227,0.0005353154,0.007336294,0.8983351,0.06287038],"study_design_scores_gemma":[0.0007515399,0.0006474633,0.0355086,0.0008862403,0.0002652321,0.0007833531,0.03466606,0.003571612,0.0006079875,0.01927167,0.9026954,0.0003448486],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.4444601,0.2176493,0.001174416,0.1974384,0.007613519,0.00304316,0.0000609031,0.0000528602,0.1285074],"genre_scores_gemma":[0.912503,0.01561931,0.01753966,0.03890968,0.0009196532,0.0001867853,0.00001330005,0.00003999101,0.01426864],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.4680429,"threshold_uncertainty_score":0.9986572,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.3834619295578525,"score_gpt":0.5385322126661344,"score_spread":0.1550702831082819,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}