{"id":"W309965191","doi":"","title":"Alternative Strategies for Large Scale Student Assessment in Canada: Is Value-Added Assessment One Possible Answer","year":2005,"lang":"en","type":"article","venue":"Canadian Journal of Educational Administration and Policy","topic":"Evaluation and Performance Assessment","field":"Decision Sciences","cited_by":21,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Accountability; Scale (ratio); Popularity; Value (mathematics); Educational assessment; Student achievement; Public relations; Academic achievement; Political science; Psychology; Pedagogy; Computer science; Geography; Social psychology","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":true,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.04157157,0.0009199967,0.001079702,0.002911492,0.004056776,0.009071456,0.00475527,0.003981158,0.004798912],"category_scores_gemma":[0.1327812,0.0005246936,0.0007781197,0.006057275,0.006438337,0.007508584,0.004775245,0.003855069,0.0005905571],"about_ca_system_candidate":true,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.04549113,"about_ca_system_score_gemma":0.1157444,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_topic_score_codex":0.8547796,"about_ca_topic_score_gemma":0.892382,"domain_scores_codex":[0.9596454,0.02146253,0.001289537,0.002324331,0.01212137,0.003156856],"domain_scores_gemma":[0.8878789,0.04552132,0.007682887,0.007553187,0.04527216,0.00609155],"domain_codex":null,"domain_gemma":"evaluation","domain_candidate":"evaluation","domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.000293098,0.0003324855,0.03254082,0.001076607,0.0002859173,0.0003397131,0.005706951,0.02408358,0.0005969838,0.5286959,0.06532312,0.3407248],"study_design_scores_gemma":[0.0006019665,0.0003864334,0.05465543,0.002577342,0.0001950906,0.0003577192,0.01564489,0.1438947,0.001764984,0.6105054,0.1686918,0.0007242159],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"commentary","genre_gemma":"empirical","genre_scores_codex":[0.1427843,0.01059233,0.1652958,0.5065297,0.001197667,0.001973743,0.001312699,0.001169284,0.1691444],"genre_scores_gemma":[0.8750351,0.002335663,0.1002467,0.009332824,0.0002366073,0.000617592,0.0002439293,0.00006921423,0.01188241],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.9584284,"threshold_uncertainty_score":0.3300628,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.1030903987976077,"score_gpt":0.4976632851093496,"score_spread":0.3945728863117419,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}