{"id":"W2061917416","doi":"10.1111/emip.12015","title":"The Multiple‐Use of Accountability Assessments: Implications for the Process of Validation","year":2013,"lang":"en","type":"article","venue":"Educational Measurement Issues and Practice","topic":"Educational Assessment and Improvement","field":"Decision Sciences","cited_by":16,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Manitoba","funders":"","keywords":"Accountability; Process (computing); Argument (complex analysis); Quality (philosophy); Management science; Process management; Computer science; Best practice; Psychology; Political science; Medicine; Business; Engineering; Epistemology","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.721591,0.001751917,0.003008878,0.01031651,0.0254986,0.03246795,0.009584797,0.01394565,0.002968446],"category_scores_gemma":[0.7751388,0.002446288,0.00257444,0.009981265,0.1234637,0.04980734,0.02784471,0.02041711,0.0007283381],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.03987854,"about_ca_system_score_gemma":0.09643989,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.02375266,"about_ca_topic_score_gemma":0.02391216,"domain_scores_codex":[0.145104,0.7647716,0.02317503,0.01155008,0.04986543,0.00553385],"domain_scores_gemma":[0.0600758,0.8492451,0.02181414,0.03232497,0.03292128,0.003618742],"domain_codex":"methods","domain_gemma":"evaluation","domain_candidate":"evaluation","domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.00008476834,0.0001642141,0.01158504,0.0006922293,0.00007893084,0.000557445,0.1245465,0.001250456,0.0003783683,0.797132,0.002924558,0.0606056],"study_design_scores_gemma":[0.0001170672,0.0002251781,0.007423201,0.005104395,0.000046848,0.0007739118,0.06948519,0.007897135,0.001236501,0.8553298,0.05201752,0.0003433668],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.06633445,0.007421612,0.5709605,0.2760496,0.001408725,0.003918594,0.00009930593,0.000401552,0.07340566],"genre_scores_gemma":[0.7337202,0.001123954,0.2519969,0.006191339,0.0002988302,0.003724456,0.00003307662,0.0001414417,0.002769921],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.278409,"threshold_uncertainty_score":0.3433279,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.4021609414534796,"score_gpt":0.5386433098125486,"score_spread":0.136482368359069,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}