{"id":"W4253831379","doi":"10.1017/s0261444806243313","title":"Language testing","year":2006,"lang":"en","type":"article","venue":"Language Teaching","topic":"Second Language Learning and Teaching","field":"Arts and Humanities","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Grading (engineering); Literacy; Vocabulary; Psychology; Sociology; Pedagogy; Philosophy; Engineering; Linguistics","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.003437908,0.0009265826,0.0007015497,0.002071712,0.001449349,0.003666956,0.001825823,0.001573426,0.4062755],"category_scores_gemma":[0.01156318,0.0003757211,0.0007864135,0.001408572,0.0008949857,0.002764321,0.004148727,0.002041174,0.3304473],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001418901,"about_ca_system_score_gemma":0.003098212,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.003130326,"about_ca_topic_score_gemma":0.007576908,"domain_scores_codex":[0.9963346,0.0007991183,0.0003596164,0.0004957296,0.001718925,0.0002921215],"domain_scores_gemma":[0.9924545,0.0009695183,0.0002258117,0.001389754,0.004075247,0.000884964],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","study_design_scores_codex":[0.0001015579,0.0001097031,0.001623595,0.0001365142,0.000009587141,0.0002799224,0.0004559826,0.00006644698,0.0006119338,0.008109296,0.5467669,0.4417286],"study_design_scores_gemma":[0.00001511737,0.00006741125,0.001833623,0.0001825056,0.000003770191,0.0004677327,0.0003256885,0.0000825895,0.0005077876,0.002853637,0.9936474,0.00001263278],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","genre_codex":"other","genre_gemma":"empirical","genre_scores_codex":[0.003948737,0.002206397,0.006107792,0.007528793,0.003748569,0.0004348487,0.006139614,0.003843621,0.9660416],"genre_scores_gemma":[0.02261921,0.00176987,0.009035088,0.004224514,0.0005284852,0.0004556457,0.006547328,0.001321683,0.9534983],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.4062755,"threshold_uncertainty_score":0.846876,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.01691662861630113,"score_gpt":0.2350651121370518,"score_spread":0.2181484835207507,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}