{"id":"W2384269281","doi":"","title":"A Study on Canadian Language Benchmarks","year":2003,"lang":"en","type":"article","venue":"Journal of Hunan University","topic":"Higher Education and Teaching Methods","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Test of English as a Foreign Language; Computer science; Task (project management); Language assessment; Linguistics; Language proficiency; Scale (ratio); Mathematics education; Natural language processing; Psychology; Engineering; Geography","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004934519,0.0006637356,0.0007281799,0.005090812,0.02040707,0.004426755,0.002065601,0.0009545975,0.006363511],"category_scores_gemma":[0.02482246,0.0005509045,0.0004480857,0.01476182,0.002861628,0.002005548,0.003020002,0.001960617,0.0006203254],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.08383507,"about_ca_system_score_gemma":0.1034149,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_topic_score_codex":0.9917788,"about_ca_topic_score_gemma":0.9952757,"domain_scores_codex":[0.9923372,0.001032689,0.0002803478,0.0005851509,0.003874679,0.001889946],"domain_scores_gemma":[0.9792712,0.00285562,0.001069959,0.0006156223,0.01340622,0.002781374],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"qualitative","study_design_gemma":"not_applicable","study_design_scores_codex":[0.0005573597,0.001420773,0.2499343,0.0005946248,0.00006333767,0.003023817,0.530695,0.0005489419,0.00146076,0.02993646,0.05434636,0.1274183],"study_design_scores_gemma":[0.00005024971,0.0003921564,0.345191,0.0003227314,0.00003124736,0.0004353994,0.5265403,0.0005219035,0.0006298136,0.0007179563,0.1249998,0.0001674657],"study_design_candidate":"not_applicable","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9472014,0.0006646073,0.0001372959,0.001947238,0.00009153191,0.0002072215,0.001178725,0.00001549895,0.04855644],"genre_scores_gemma":[0.9846669,0.0009751465,0.0003999276,0.0007428945,0.00001683265,0.0001715973,0.001348541,0.0000356524,0.01164253],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.08383507,"threshold_uncertainty_score":0.6082689,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.01883042602821232,"score_gpt":0.2813370304773516,"score_spread":0.2625066044491393,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}