{"id":"W2376153660","doi":"","title":"The Benchmarks for Reading of Canadian Language Benchmarks and its Comments","year":2014,"lang":"en","type":"article","venue":"Overseas English","topic":"Second Language Learning and Teaching","field":"Arts and Humanities","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Scope (computer science); Reading (process); Immigration; Strengths and weaknesses; Computer science; Linguistics; Language assessment; Psychology; Mathematics education; Political science; Social psychology","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0004210596,0.00008486577,0.0001076529,0.0001027533,0.0005358445,0.0001325023,0.0001105196,0.00003484748,0.001025341],"category_scores_gemma":[0.0006114388,0.00006456004,0.00004431773,0.00002466442,0.00005544979,0.00009842247,0.00002018599,0.0001338491,5.805679e-7],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.00002046729,"about_ca_system_score_gemma":0.00002056405,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_topic_score_codex":0.03639047,"about_ca_topic_score_gemma":0.1225546,"domain_scores_codex":[0.9994044,0.00005548998,0.0001156961,0.0001178495,0.00007855789,0.0002279713],"domain_scores_gemma":[0.9992601,0.0003775786,0.0000638637,0.0001308154,0.00007176615,0.00009593688],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","study_design_scores_codex":[0.00001430709,0.00001072513,0.000477954,0.000049988,0.00005447835,0.000001696342,0.1430391,0.000005253259,0.00002733528,0.8107374,0.02849056,0.01709121],"study_design_scores_gemma":[0.0003047545,0.00004197529,0.0001400941,0.00003142305,0.00002010635,4.317942e-7,0.0417595,0.0005219303,0.0000412389,0.0001007348,0.9569443,0.00009348834],"study_design_candidate":"not_applicable","study_design_consensus":null,"genre_codex":"other","genre_gemma":"empirical","genre_scores_codex":[0.2497992,0.0004369438,0.000001977744,0.0001208082,0.000601299,0.0001400576,0.00005697693,0.00001912562,0.7488236],"genre_scores_gemma":[0.9959094,0.00001345545,0.00002168464,0.0003592667,0.0005622311,0.000009586412,0.00009295355,0.00001191889,0.003019522],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.9284537,"threshold_uncertainty_score":0.9998879,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.01115129890597977,"score_gpt":0.2094686924585338,"score_spread":0.1983173935525541,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}