{"id":"W2088378658","doi":"10.1002/pits.20019","title":"The early assessment conundrum: Lessons from the past, implications for the future","year":2004,"lang":"en","type":"article","venue":"Psychology in the Schools","topic":"Educational and Psychological Assessments","field":"Psychology","cited_by":33,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Toronto","funders":"U.S. Department of Education","keywords":"Psychology; Accountability; Context (archaeology); Early childhood education; Early childhood; Educational assessment; Developmental psychology; Pedagogy; Political science","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.1112684,0.00130088,0.001585104,0.004821123,0.004483302,0.009683729,0.007185596,0.01269239,0.005051855],"category_scores_gemma":[0.1031695,0.0007101727,0.001052376,0.003216217,0.01962729,0.02228541,0.007937067,0.02210594,0.00105226],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.01266889,"about_ca_system_score_gemma":0.02544129,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.01617354,"about_ca_topic_score_gemma":0.03525234,"domain_scores_codex":[0.9688314,0.01955978,0.003394129,0.001804392,0.004785969,0.001624371],"domain_scores_gemma":[0.8786004,0.06363723,0.004826344,0.006885737,0.03014939,0.01590108],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.0003226313,0.0008139412,0.009884111,0.003626755,0.00005914908,0.002259857,0.005467678,0.001123191,0.0002900398,0.1172228,0.1512922,0.7076377],"study_design_scores_gemma":[0.0002132446,0.0006214405,0.01983957,0.03851525,0.0001204817,0.01025497,0.05472817,0.005387228,0.0008507628,0.4155671,0.4534071,0.0004946485],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"genre_codex":"commentary","genre_gemma":"empirical","genre_scores_codex":[0.002836344,0.1115474,0.004857685,0.8721822,0.004525749,0.00005417382,0.0000502655,0.00009717949,0.003849108],"genre_scores_gemma":[0.3266451,0.3514755,0.101367,0.1930809,0.01773513,0.001071843,0.0001822218,0.0002028867,0.008239497],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.1112684,"threshold_uncertainty_score":0.5884507,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.0794191862697233,"score_gpt":0.4694691489686061,"score_spread":0.3900499626988828,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}