{"id":"W98146743","doi":"10.1007/978-94-007-5902-2_16","title":"Authentic Assessment, Teacher Judgment and Moderation in a Context of High Accountability","year":2014,"lang":"en","type":"book-chapter","venue":"The enabling power of assessment","topic":"Student Assessment and Feedback","field":"Social Sciences","cited_by":11,"is_retracted":false,"has_abstract":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Accountability; Moderation; Context (archaeology); Psychology; Political science; Social psychology; History; Law; Archaeology","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaepi_narrow"],"consensus_categories":[],"category_scores_codex":[0.003470941,0.0004250066,0.0009102262,0.0002778387,0.0002796599,0.0001579254,0.0005469701,0.0003903435,0.000777158],"category_scores_gemma":[0.00002572135,0.0003561444,0.0002103186,0.0001100841,0.0006266482,0.0001902668,0.0002431199,0.0006233383,0.000005562828],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0005895839,"about_ca_system_score_gemma":0.000637068,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.003447302,"about_ca_topic_score_gemma":0.005740148,"domain_scores_codex":[0.9961929,0.0002848635,0.001058683,0.0006111035,0.001430097,0.0004223809],"domain_scores_gemma":[0.997518,0.0004493717,0.0009984577,0.0006196297,0.0003132479,0.000101278],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","study_design_scores_codex":[0.0001505319,0.0008335209,0.02901954,0.0004140405,0.0007521383,0.000005833489,0.03089238,0.00009448195,0.0004079991,0.9210027,0.001444462,0.01498239],"study_design_scores_gemma":[0.01700641,0.003290631,0.2640035,0.004862257,0.002991076,0.000006771932,0.130587,0.007001617,0.0001909516,0.2315556,0.3321404,0.006363742],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"genre_codex":"other","genre_gemma":"empirical","genre_scores_codex":[0.3174408,0.0008931434,0.005757918,0.003075094,0.002358162,0.005043742,0.00008873788,0.0001241446,0.6652183],"genre_scores_gemma":[0.9579663,0.0002320219,0.0004891652,0.0001131073,0.0001314654,0.00005064488,0.00003093699,0.00004596591,0.04094042],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.6894471,"threshold_uncertainty_score":0.9998891,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.03051451189430561,"score_gpt":0.3371362410140897,"score_spread":0.3066217291197841,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}