{"id":"W1799084732","doi":"10.3386/w16877","title":"The Effect of Evaluation on Performance: Evidence from Longitudinal Student Achievement Data of Mid-career Teachers","year":2011,"lang":"en","type":"report","venue":"National Bureau of Economic Research","topic":"School Choice and Performance","field":"Social Sciences","cited_by":63,"is_retracted":false,"has_abstract":true,"ca_institutions":"Manning Diversified Forest Products (Canada)","funders":"Institute of Education Sciences; Harvard University; Joyce Foundation; U.S. Department of Education","keywords":"Mathematics education; Longitudinal data; Student achievement; Psychology; Academic achievement; Econometrics; Medical education; Computer science; Mathematics; Medicine; Data mining","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008869334,0.0002434382,0.0005060471,0.001381688,0.0006595266,0.00133366,0.0009025735,0.0009905201,0.00252909],"category_scores_gemma":[0.03697263,0.0003786081,0.0006335685,0.001970976,0.0005926772,0.0009924273,0.00137428,0.001465183,0.001141596],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0006119227,"about_ca_system_score_gemma":0.0009657554,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.02350374,"about_ca_topic_score_gemma":0.02994801,"domain_scores_codex":[0.9951741,0.002455732,0.0004249916,0.0005843233,0.000819956,0.0005409591],"domain_scores_gemma":[0.9019832,0.04067127,0.03654321,0.007622966,0.00704462,0.006134818],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"observational","study_design_scores_codex":[0.0001546197,0.0002614113,0.9928784,0.00001278122,0.00008879044,0.0000277843,0.0004743922,0.0002066433,0.00007746492,0.00009581602,0.0003081264,0.005413817],"study_design_scores_gemma":[0.000006131553,0.0001265694,0.9988052,0.00001113091,0.00002523918,0.00001587723,0.0002144868,0.0002884479,0.00008778426,0.00007927584,0.0003348608,0.000005117152],"study_design_candidate":"observational","study_design_consensus":"observational","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.997681,0.0003603302,0.0002223455,0.0002522485,0.000008117692,0.000007415099,0.0006694504,0.000008215567,0.0007909251],"genre_scores_gemma":[0.9976097,0.0002468523,0.0001275839,0.00003726053,0.00001435987,0.00001571991,0.001097463,0.000005715798,0.0008453109],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.02350374,"threshold_uncertainty_score":0.04690611,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.6749455625049929,"score_gpt":0.6025124768252421,"score_spread":0.07243308567975082,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}