{"id":"W4391036586","doi":"10.15864/ijelts.6207","title":"IELTS Recent Developments Examined for their Impact and Effectiveness for Evaluating Language Proficiency of Candidates","year":2024,"lang":"en","type":"article","venue":"International Journal of English Learning & Teaching Skills","topic":"Multilingual Education and Policy","field":"Social Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Computer science; Linguistics; Natural language processing; Psychology; Language proficiency; Mathematics education; Philosophy","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03103796,0.0005043881,0.0005819366,0.005590052,0.000710836,0.002523484,0.001246956,0.0009020616,0.008720547],"category_scores_gemma":[0.07208081,0.0002937563,0.001041949,0.005321728,0.0008822146,0.003525314,0.001531419,0.0018124,0.002138526],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.00337049,"about_ca_system_score_gemma":0.003782979,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.008865378,"about_ca_topic_score_gemma":0.01348843,"domain_scores_codex":[0.9844316,0.006711693,0.001775922,0.0009512784,0.005534507,0.0005950672],"domain_scores_gemma":[0.9204195,0.02925068,0.0103758,0.0022303,0.03450726,0.00321642],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"observational","study_design_scores_codex":[0.001392002,0.000940633,0.1239381,0.001586672,0.0001208554,0.0001673925,0.0008916899,0.0005874619,0.001024678,0.003882597,0.0217452,0.8437228],"study_design_scores_gemma":[0.000351818,0.006197318,0.7470095,0.004097454,0.0006115935,0.001365485,0.00429848,0.003763876,0.008048796,0.002066504,0.2219159,0.0002733678],"study_design_candidate":"observational","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.622662,0.06189119,0.02321855,0.03909451,0.004684024,0.003604389,0.01065841,0.001973381,0.2322135],"genre_scores_gemma":[0.8310512,0.03574076,0.06139524,0.004986932,0.001786594,0.002381459,0.00829875,0.0002937449,0.05406528],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.03103796,"threshold_uncertainty_score":0.1641464,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.03447392864129107,"score_gpt":0.493671357061322,"score_spread":0.4591974284200309,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}