{"id":"W438995078","doi":"","title":"At Odds: School Achievement -- Bad Data Must Be Challenged","year":2000,"lang":"en","type":"article","venue":"Phi Delta Kappan","topic":"Educational and Psychological Assessments","field":"Psychology","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Surprise; Credibility; Odds; Psychology; Statement (logic); Mathematics education; Law; Sociology; Social psychology; Political science; Computer science","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0003979204,0.0002510414,0.0002363204,0.00006178443,0.0002215887,0.00003934502,0.001100252,0.0001537642,0.1259904],"category_scores_gemma":[0.0000267586,0.0002123168,0.00006767632,0.0002000219,0.00009276574,0.00016161,0.0002170501,0.0003153916,0.009854015],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.00007799945,"about_ca_system_score_gemma":0.00003565918,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.0001632484,"about_ca_topic_score_gemma":0.000108736,"domain_scores_codex":[0.9976941,0.0001814682,0.000387735,0.0008717445,0.0003536094,0.0005113371],"domain_scores_gemma":[0.9978748,0.0001193238,0.0000863399,0.001556835,0.00004036356,0.0003223279],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","study_design_scores_codex":[0.0006687439,0.005599367,0.02859136,0.0000323444,0.0006467138,0.00007252458,0.002687159,0.00001348028,0.0002560492,0.01175501,0.8062854,0.1433919],"study_design_scores_gemma":[0.001004559,0.0001475587,0.2250115,0.0000136442,0.00003560473,0.00001874178,0.0003938434,0.00002841388,0.000005702864,0.0009655436,0.7720827,0.0002922572],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.7010319,0.0008742007,0.00003151635,0.02138697,0.0009713775,0.0004125529,0.0003974183,0.0001392579,0.2747548],"genre_scores_gemma":[0.9371834,0.0002398407,0.000414214,0.007495831,0.0008873233,0.000152495,0.001658511,0.00003597437,0.05193238],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.2361515,"threshold_uncertainty_score":0.9909169,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.1811150590644432,"score_gpt":0.4199190268113192,"score_spread":0.238803967746876,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}