{"id":"W220468415","doi":"","title":"THE PARADOX OF CLASSROOM ASSESSMENT: A CHALLENGE FOR THE 21ST CENTURY","year":2001,"lang":"en","type":"article","venue":"McGill Journal of Education / Revue des sciences de l'éducation de McGill","topic":"Education and Critical Thinking Development","field":"Social Sciences","cited_by":7,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Toronto","funders":"","keywords":"Competence (human resources); Humanities; Curriculum; Sociology; Ethnology; Philosophy; Pedagogy; Psychology; Social psychology","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.006605851,0.0001286805,0.0001831618,0.0001697764,0.004681317,0.0001517062,0.0009779646,0.00008570731,0.0001188161],"category_scores_gemma":[0.00239177,0.00008790012,0.000167594,0.0008179005,0.001146917,0.0004054336,0.00002966319,0.000230464,0.000005344941],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001242488,"about_ca_system_score_gemma":0.002461502,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.0001551867,"about_ca_topic_score_gemma":0.0004405637,"domain_scores_codex":[0.9976236,0.0004681459,0.0006816656,0.0001951729,0.0005736022,0.0004578433],"domain_scores_gemma":[0.9957335,0.001922537,0.0006883713,0.0002251841,0.001204637,0.0002257898],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","study_design_scores_codex":[0.00001938729,0.000500195,0.00127481,0.00003747896,0.00003629603,3.719097e-7,0.01344989,0.00008535079,0.00004828271,0.8693834,0.002536842,0.1126277],"study_design_scores_gemma":[0.0001418815,0.0001252494,0.004156366,0.0001225793,0.00004731707,0.00005903991,0.09072834,0.0001088992,0.00005624296,0.0997512,0.8045824,0.0001204761],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"genre_codex":"commentary","genre_gemma":"empirical","genre_scores_codex":[0.1442328,0.01133499,0.006857294,0.4558126,0.02746417,0.002602464,0.00006183556,0.0001037068,0.3515301],"genre_scores_gemma":[0.9659639,0.01446787,0.01541564,0.00174476,0.0006140058,0.00008910471,0.000002131606,0.00001356107,0.001689014],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.8217311,"threshold_uncertainty_score":0.9966145,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.2302648335314056,"score_gpt":0.4522286731329828,"score_spread":0.2219638396015772,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}