{"id":"W2038816208","doi":"10.5539/hes.v1n2p107","title":"Fairness of IELTS Test Scores in University Admission","year":2011,"lang":"en","type":"article","venue":"Higher Education Studies","topic":"Student Assessment and Feedback","field":"Social Sciences","cited_by":1,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Test (biology); Context (archaeology); Language proficiency; Psychology; Language assessment; Interpretation (philosophy); Empirical research; Medical education; Mathematics education; Applied psychology; Medicine; Computer science; Statistics","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":true,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0001487387,0.00004934005,0.0001021431,0.00008653036,0.0001322572,0.00000464505,0.0001203217,0.00002973927,0.0004102839],"category_scores_gemma":[0.00003474929,0.00004643317,0.00002139868,0.0003176765,0.0001534819,0.0001348145,0.00004133474,0.00003332619,0.00001592353],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.00007590156,"about_ca_system_score_gemma":0.0001652939,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001534164,"about_ca_topic_score_gemma":0.0005882817,"domain_scores_codex":[0.9995089,0.00005386213,0.00009200403,0.0001022193,0.0001444982,0.00009849593],"domain_scores_gemma":[0.999559,0.0001018927,0.00006785239,0.00006790651,0.0001707041,0.00003257397],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"observational","study_design_scores_codex":[0.000008872671,0.0003596725,0.9336359,0.00001187904,0.0000127465,4.511861e-7,0.03114586,3.474394e-8,0.00002638799,0.02926815,0.005210741,0.0003193259],"study_design_scores_gemma":[0.0001089289,0.00001573436,0.8861148,0.00004007128,0.000008664788,1.611469e-8,0.08407332,3.817646e-8,0.00005440713,0.0007171545,0.02881438,0.00005250534],"study_design_candidate":"observational","study_design_consensus":"observational","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.8835381,0.0006869542,9.733327e-7,0.0005637364,0.0008405375,0.00009772788,7.953055e-7,0.00001950674,0.1142517],"genre_scores_gemma":[0.9770831,0.000452649,0.0001808083,0.0000260594,0.00007471812,0.000002805561,7.053529e-7,0.000002340982,0.0221768],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.09354506,"threshold_uncertainty_score":0.4492322,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.1157244094817193,"score_gpt":0.3958591156896718,"score_spread":0.2801347062079525,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}