{"id":"W4254740339","doi":"10.1017/s0261444805242398","title":"Language testing","year":2004,"lang":"en","type":"article","venue":"Language Teaching","topic":"Second Language Learning and Teaching","field":"Arts and Humanities","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"Purdue University","keywords":"Formative assessment; Language assessment; German; Psychology; Vocabulary; Comparability; Competence (human resources); Pedagogy; Library science; History; Linguistics; Computer science; Mathematics; Social psychology","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.008213534,0.001248426,0.0009342649,0.00295555,0.001998076,0.003921309,0.00251016,0.001681315,0.4134638],"category_scores_gemma":[0.03113158,0.0006105208,0.001072831,0.002211708,0.001130012,0.003045032,0.00491379,0.00292927,0.2964275],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.002486045,"about_ca_system_score_gemma":0.007071311,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.008120982,"about_ca_topic_score_gemma":0.01386708,"domain_scores_codex":[0.9909876,0.002338123,0.001037547,0.001173699,0.003862029,0.0006009694],"domain_scores_gemma":[0.9785492,0.002662041,0.0006703185,0.003230265,0.0130284,0.001859825],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","study_design_scores_codex":[0.0002221597,0.0002171243,0.00291873,0.0002369623,0.00001637806,0.0002978443,0.0009215236,0.00009897693,0.0006942095,0.006163508,0.6466525,0.34156],"study_design_scores_gemma":[0.00004807522,0.0001610689,0.004993424,0.0004889001,0.000009501139,0.0005564151,0.0007830488,0.0001771017,0.000620322,0.003282838,0.9888509,0.00002835132],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","genre_codex":"other","genre_gemma":"other","genre_scores_codex":[0.007700022,0.00226639,0.01264134,0.009699746,0.004012933,0.001984491,0.02069917,0.006710864,0.934285],"genre_scores_gemma":[0.03504293,0.00240002,0.01594554,0.006285274,0.000518919,0.00242168,0.01851629,0.00225421,0.9166151],"genre_candidate":"other","genre_consensus":"other","teacher_disagreement_score":0.4134638,"threshold_uncertainty_score":0.8366227,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.02446946919005953,"score_gpt":0.2524834481474602,"score_spread":0.2280139789574007,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}