{"id":"W2945229784","doi":"10.5539/ijel.v9n3p319","title":"Introducing and Testing a Measurement Tool for English Language Proficiency: Aisha’s Tool","year":2019,"lang":"en","type":"article","venue":"International Journal of English Linguistics","topic":"Reading and Literacy Development","field":"Psychology","cited_by":1,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Test (biology); Sentence; Language assessment; Test of English as a Foreign Language; Psychology; English language; Affect (linguistics); Language proficiency; Mathematics education; Natural language processing; Computer science; Communication","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":true,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004587327,0.0004186512,0.0004571795,0.004328357,0.000781321,0.00120575,0.0006561277,0.0006789387,0.006017487],"category_scores_gemma":[0.01918451,0.0002668478,0.0005863276,0.00142311,0.0006671558,0.001554775,0.001579119,0.0008951506,0.002446158],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0004113357,"about_ca_system_score_gemma":0.00167087,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001187775,"about_ca_topic_score_gemma":0.002242213,"domain_scores_codex":[0.9947037,0.00185907,0.000807026,0.0003466049,0.001991459,0.0002920032],"domain_scores_gemma":[0.9824902,0.009116298,0.001512739,0.001085155,0.005146527,0.0006490145],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"observational","study_design_scores_codex":[0.0004003604,0.002059251,0.4238344,0.0005265122,0.0001083187,0.000399362,0.007151698,0.001237672,0.01263154,0.005902748,0.02052942,0.5252187],"study_design_scores_gemma":[0.0001704399,0.003430614,0.8693326,0.0005319911,0.0001292779,0.00131332,0.01036186,0.01060805,0.02239071,0.005024511,0.07646471,0.0002418431],"study_design_candidate":"observational","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.8606864,0.0003995743,0.04167621,0.001248413,0.000577185,0.003949022,0.003324392,0.00190783,0.08623104],"genre_scores_gemma":[0.7949989,0.0004737953,0.1780254,0.0004922384,0.0001118083,0.006075219,0.001521198,0.00009088483,0.01821062],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.006017487,"threshold_uncertainty_score":0.02426034,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.02417356471455197,"score_gpt":0.3095135211335783,"score_spread":0.2853399564190264,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}