{"id":"W146577395","doi":"","title":"How Do Other Countries Evaluate Teachers? Countries Known to Outpace the U.S. in Student Achievement Use a Variety of Educational and Organizational Methods, but Rarely Use the Approaches to Education Reform That We Are Promoting","year":2012,"lang":"en","type":"article","venue":"Phi Delta Kappan","topic":"Teacher Education and Leadership Studies","field":"Social Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Accountability; Variety (cybernetics); Academic achievement; Test (biology); Standardized test; Mathematics education; Scale (ratio); Achievement test; Teacher education; Educational assessment; Pedagogy; Political science; Psychology; Public relations; Computer science","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.02597616,0.000641189,0.001247381,0.004675623,0.005995463,0.01515883,0.001247338,0.001914467,0.00559525],"category_scores_gemma":[0.07114781,0.0005270602,0.0006980005,0.009771632,0.004658682,0.01159005,0.003414291,0.002920374,0.002499961],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.01304685,"about_ca_system_score_gemma":0.01578904,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.2173285,"about_ca_topic_score_gemma":0.187642,"domain_scores_codex":[0.967679,0.01809033,0.002554779,0.002547655,0.005892766,0.00323542],"domain_scores_gemma":[0.9068515,0.01203519,0.01069115,0.008672532,0.05314421,0.008605397],"domain_codex":null,"domain_gemma":"evaluation","domain_candidate":"evaluation","domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"observational","study_design_scores_codex":[0.0002822836,0.0002681016,0.5062631,0.001119605,0.0005849389,0.0003814667,0.04542101,0.001205032,0.001063982,0.03009821,0.1124191,0.3008931],"study_design_scores_gemma":[0.0001487186,0.0003520756,0.4830329,0.004930383,0.000469873,0.0003885331,0.1565536,0.001253726,0.003371561,0.01216145,0.3369098,0.0004272426],"study_design_candidate":"observational","study_design_consensus":"observational","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.4564652,0.02967515,0.01822253,0.1634365,0.004970418,0.001077106,0.006141643,0.000625519,0.3193859],"genre_scores_gemma":[0.9258304,0.007828137,0.01594116,0.03075457,0.0003776625,0.0005474202,0.002098642,0.0004488383,0.01617323],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.9740238,"threshold_uncertainty_score":0.4321271,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.2937060094397593,"score_gpt":0.4087077009401237,"score_spread":0.1150016915003644,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}