{"id":"W3144739624","doi":"","title":"Test Review： Canadian Academic English Language （CAEL） Assessment","year":2011,"lang":"en","type":"article","venue":"Acta Scientiarum Naturalium Universitatis Sunyatseni","topic":"Educational Technology and Assessment","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"ca_institutions":"","funders":"","keywords":"Test (biology); Language assessment; Linguistics; Computer science; Natural language processing; Geology; Philosophy","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000541489,0.0002388335,0.0002266933,0.0005815215,0.000459579,0.00006993652,0.002205483,0.0002052504,0.0004507067],"category_scores_gemma":[0.0002308519,0.0002368699,0.0001053733,0.001521517,0.0001977966,0.001501981,0.0003889143,0.0007610368,0.0001009467],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0004222179,"about_ca_system_score_gemma":0.0009254398,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.004952585,"about_ca_topic_score_gemma":0.00542621,"domain_scores_codex":[0.9978423,0.0001020528,0.0002636206,0.0006595814,0.0004703831,0.0006620429],"domain_scores_gemma":[0.9981379,0.0001836823,0.0001758975,0.0007971459,0.0003129215,0.0003924523],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","study_design_scores_codex":[0.000005587678,0.0002999245,0.03558251,0.000201706,0.000167707,0.0004182275,0.0134499,0.000001411054,0.004326071,0.4144393,0.5216288,0.009478875],"study_design_scores_gemma":[0.003073712,0.0009494256,0.4406343,0.002420126,0.0005289488,0.0003846269,0.02498738,0.002820311,0.01390611,0.01400471,0.4913099,0.004980414],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","genre_codex":"other","genre_gemma":"empirical","genre_scores_codex":[0.3456495,0.02028182,0.01004469,0.1771273,0.02217707,0.004038028,0.0004189095,0.003508229,0.4167545],"genre_scores_gemma":[0.9622161,0.000585012,0.03243553,0.001876767,0.00005917979,0.000008072756,0.00005050128,0.00001158666,0.002757221],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.6165666,"threshold_uncertainty_score":0.9659273,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.01311408641835332,"score_gpt":0.2655552543664973,"score_spread":0.252441167948144,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}