{"id":"W276684465","doi":"","title":"The One-Legged High Jumper and the Perils of Prediction: Predicting Success for Students Based on Their Background Is More Accurate in the Aggregate Than in Individual Situations, Where It Should Never Be Applied","year":2012,"lang":"en","type":"article","venue":"Phi Delta Kappan","topic":"Reading and Literacy Development","field":"Psychology","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Surprise; Psychology; Reading (process); Graduation (instrument); Mathematics education; Poverty; Literacy; Social psychology; Pedagogy; Political science; Mathematics","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.03172676,0.001361727,0.002044328,0.004658969,0.002130673,0.006566034,0.002100597,0.003120109,0.002659258],"category_scores_gemma":[0.1071254,0.0006655161,0.001286184,0.003817925,0.004641998,0.008252067,0.004047072,0.006514825,0.001160632],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001969111,"about_ca_system_score_gemma":0.002579215,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.03752236,"about_ca_topic_score_gemma":0.02128959,"domain_scores_codex":[0.9773037,0.01383222,0.0009919845,0.003211156,0.004099985,0.000560914],"domain_scores_gemma":[0.8735208,0.09320398,0.01260397,0.007467123,0.006824943,0.006379247],"domain_codex":null,"domain_gemma":"methods","domain_candidate":"methods","domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.0003206216,0.0002518568,0.9088664,0.000178432,0.0007929287,0.00008061136,0.001338139,0.002148666,0.00006462022,0.003387795,0.01008247,0.07248745],"study_design_scores_gemma":[0.0001128696,0.001218534,0.8230327,0.001707352,0.000632903,0.000455032,0.004858936,0.05900751,0.0004066174,0.09771033,0.01050994,0.0003472433],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.7869275,0.02093894,0.04426691,0.1183271,0.002753636,0.0003338864,0.003343524,0.0005969797,0.02251159],"genre_scores_gemma":[0.9842979,0.001954828,0.006699387,0.004199979,0.0008853439,0.0001137749,0.0007414619,0.00003828116,0.001069009],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.9682732,"threshold_uncertainty_score":0.1677892,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.08755749517404587,"score_gpt":0.3703974535491772,"score_spread":0.2828399583751313,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}