{"id":"W1948718713","doi":"10.1007/s11092-015-9231-8","title":"Towards a framework for the validation of early childhood assessment systems","year":2015,"lang":"en","type":"article","venue":"Educational Assessment Evaluation and Accountability","topic":"Early Childhood Education and Development","field":"Social Sciences","cited_by":47,"is_retracted":false,"has_abstract":false,"ca_institutions":"York University","funders":"Fok Ying Tong Education Foundation; U.S. Department of Education","keywords":"Early childhood education; Early childhood; Accountability; Educational assessment; Conceptual framework; Argument (complex analysis); Field (mathematics); Work (physics); Psychology; Political science; Pedagogy; Developmental psychology; Sociology; Medicine; Social science; Engineering","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.4722507,0.001904739,0.00259622,0.01264423,0.006790054,0.0266776,0.0113544,0.008099965,0.002004541],"category_scores_gemma":[0.5449938,0.001891721,0.002673602,0.006368055,0.01755885,0.01755909,0.01578632,0.0107188,0.0007505732],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.02485839,"about_ca_system_score_gemma":0.06657612,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.04656231,"about_ca_topic_score_gemma":0.02773424,"domain_scores_codex":[0.5270795,0.3621851,0.04103117,0.01478694,0.04909934,0.005817988],"domain_scores_gemma":[0.2985397,0.3679558,0.03228508,0.09210674,0.2023206,0.006792089],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.0002301367,0.0009874721,0.04697404,0.001352689,0.0004572894,0.0002677286,0.02526402,0.01760761,0.002542695,0.7206727,0.006287778,0.1773559],"study_design_scores_gemma":[0.0003630543,0.001158953,0.04934523,0.01100633,0.0005562817,0.0004753268,0.02513688,0.1398083,0.01319085,0.6712949,0.08713847,0.0005254413],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.03009197,0.001126438,0.9260634,0.01424988,0.0002798013,0.004452029,0.0005261705,0.001002905,0.02220743],"genre_scores_gemma":[0.2081666,0.0001923042,0.7869204,0.0006414757,0.00004387689,0.0023986,0.0006888188,0.0001100297,0.0008378699],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.4722507,"threshold_uncertainty_score":0.6508089,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.09283552875712486,"score_gpt":0.4581947012178403,"score_spread":0.3653591724607154,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}