{"id":"W2606446642","doi":"","title":"The Impact of Standardised Testing on Later High Stakes Test Outcomes","year":2017,"lang":"en","type":"preprint","venue":"Lancaster EPrints (Lancaster University)","topic":"School Choice and Performance","field":"Social Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Boycott; Quarter (Canadian coin); Test (biology); Argument (complex analysis); Psychology; Educational attainment; Demographic economics; Political science; Economics; Medicine; Geography; Economic growth; Politics","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaepi_narrow"],"consensus_categories":[],"category_scores_codex":[0.000731401,0.0004394263,0.0006493652,0.0002541687,0.001224475,0.0004542331,0.001861532,0.0003575494,0.0003250154],"category_scores_gemma":[0.0005271446,0.0003159978,0.0003942173,0.0002181217,0.0005173195,0.0003584401,0.0008394549,0.0008493366,0.0002628394],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0004009134,"about_ca_system_score_gemma":0.0005044679,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00307028,"about_ca_topic_score_gemma":0.002686172,"domain_scores_codex":[0.9974207,0.0002529683,0.0003529376,0.0006046789,0.0006680108,0.0007006525],"domain_scores_gemma":[0.9961972,0.0009772653,0.0007426848,0.001466464,0.000393353,0.0002229953],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"observational","study_design_scores_codex":[0.0001876781,0.00007463559,0.9920022,0.0000487684,0.0002117649,0.00005819308,0.001660551,0.00001666697,0.00001942173,0.000217069,0.0007644072,0.004738617],"study_design_scores_gemma":[0.0007708716,0.0001243591,0.9785485,0.0003062882,0.00005571762,0.000001074184,0.000710603,0.00004937594,0.0000526128,0.0004122571,0.01858412,0.00038425],"study_design_candidate":"observational","study_design_consensus":"observational","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.8718116,0.00002914189,0.00002810503,0.0006270026,0.0007304538,0.0004573153,0.0003190701,0.00006829463,0.125929],"genre_scores_gemma":[0.974808,0.00009745849,0.000143739,0.00004418742,0.0005010187,0.000002799236,0.00001627496,0.00003306019,0.02435348],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.1029964,"threshold_uncertainty_score":0.9999292,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.06812738901110699,"score_gpt":0.3292717574372142,"score_spread":0.2611443684261072,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}