{"id":"W4244530217","doi":"10.1017/s0261444804242136","title":"Language testing","year":2004,"lang":"en","type":"article","venue":"Language Teaching","topic":"Second Language Learning and Teaching","field":"Arts and Humanities","cited_by":1,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Language assessment; Listening comprehension; Test (biology); Test of English as a Foreign Language; Interview; Psychology; Checklist; Foreign language; Library science; Active listening; Mathematics education; Sociology; Computer science","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.003298596,0.0007366335,0.0005948342,0.001784376,0.001285696,0.002595618,0.001424468,0.001071345,0.3208249],"category_scores_gemma":[0.01119498,0.000358224,0.0006238674,0.001404222,0.0006020927,0.001828799,0.00319735,0.001974182,0.235607],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001355344,"about_ca_system_score_gemma":0.003400291,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.005867755,"about_ca_topic_score_gemma":0.01227711,"domain_scores_codex":[0.9971022,0.0006266842,0.0002783909,0.0003536735,0.001367889,0.0002711511],"domain_scores_gemma":[0.9932508,0.0007186484,0.0002006693,0.001030743,0.004005553,0.0007935694],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"not_applicable","study_design_gemma":"observational","study_design_scores_codex":[0.0000815726,0.00009960325,0.001585917,0.0001394087,0.000007004966,0.0002177758,0.000481646,0.00004958266,0.0006566349,0.00366808,0.525286,0.4677268],"study_design_scores_gemma":[0.00001376132,0.00007304786,0.002810377,0.0001998264,0.000004627726,0.0004197675,0.0003581668,0.00006052076,0.0005324522,0.001614121,0.9939022,0.00001107731],"study_design_candidate":"observational","study_design_consensus":null,"genre_codex":"other","genre_gemma":"empirical","genre_scores_codex":[0.005868087,0.004006705,0.008071302,0.009099854,0.004339173,0.000643237,0.009012935,0.003664881,0.9552938],"genre_scores_gemma":[0.02502174,0.003042585,0.010176,0.005427252,0.0004803452,0.0006106232,0.008469656,0.001271926,0.9454998],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.3208249,"threshold_uncertainty_score":0.9687608,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.02446946919005953,"score_gpt":0.2524834481474602,"score_spread":0.2280139789574007,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}