{"id":"W2943008257","doi":"10.55016/ojs/jet.v40i2.52563","title":"Standardized Testing and the Classroom","year":2018,"lang":"en","type":"article","venue":"Journal of educational thought.","topic":"Student Assessment and Feedback","field":"Social Sciences","cited_by":3,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Standardized test; Psychology; Mathematics education; Pedagogy","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001777868,0.0000441939,0.0001061422,0.00004035898,0.0005146036,0.0001137707,0.0001767831,0.00002439052,0.000385305],"category_scores_gemma":[0.00201054,0.00002681684,0.00004523028,0.0002107231,0.0005504102,0.0001797227,0.00002102424,0.00009952439,0.00001125117],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.00005323654,"about_ca_system_score_gemma":0.0005609334,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00004518898,"about_ca_topic_score_gemma":0.00004152917,"domain_scores_codex":[0.9990591,0.000156048,0.0002020339,0.00005079805,0.0004231301,0.0001088972],"domain_scores_gemma":[0.9973835,0.001712522,0.0002389344,0.00004673918,0.0005571591,0.00006111919],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","study_design_scores_codex":[0.0002562993,0.00009880833,0.1496557,0.000003814284,0.0001015625,7.661051e-7,0.01458095,7.075423e-7,0.0001314922,0.7780228,0.05197497,0.00517214],"study_design_scores_gemma":[0.002833742,0.0001491193,0.2508687,0.00007974831,0.00009838322,0.00002766104,0.01500606,0.00001835494,0.00004206935,0.2088781,0.5218432,0.0001549142],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.6730266,0.0008508509,0.0001650892,0.06613798,0.002132617,0.0001712524,0.000001874428,0.000006422809,0.2575073],"genre_scores_gemma":[0.9693246,0.0001145898,0.0174973,0.0003826513,0.006862966,0.000002048635,1.904027e-7,0.000004729501,0.005810888],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.5691448,"threshold_uncertainty_score":0.4218819,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.04197803242966986,"score_gpt":0.3914393847025016,"score_spread":0.3494613522728318,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}