{"id":"W2912504907","doi":"10.22190/jtesap1803525s","title":"IS ESP STANDARDIZED ASSESSMENT FEASIBLE?","year":2019,"lang":"en","type":"article","venue":"Journal of Teaching English for Specific and Academic Purposes","topic":"EFL/ESL Teaching and Learning","field":"Arts and Humanities","cited_by":2,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Benchmarking; Standardized test; Quality assessment; Best practice; Quality assurance; Educational assessment; Language assessment; Variety (cybernetics); Quality (philosophy); Psychology; Computer science; Evaluation methods; Political science; Business; Pedagogy; External quality assessment; Engineering; Mathematics education; Operations management; Artificial intelligence; Marketing","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002372333,0.0002111606,0.0004810498,0.0002164998,0.0004712309,0.0002978416,0.0003016186,0.000117726,0.0004881678],"category_scores_gemma":[0.0002499098,0.000161339,0.0002395366,0.00002223115,0.00009088611,0.0006373316,0.00005007291,0.002034811,0.000006266327],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.00008929762,"about_ca_system_score_gemma":0.00007119124,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.000008902529,"about_ca_topic_score_gemma":0.000001129024,"domain_scores_codex":[0.9982963,0.0001459298,0.000618186,0.0002205758,0.0004142682,0.0003047668],"domain_scores_gemma":[0.9985111,0.0004032543,0.0005871259,0.0001442157,0.0002265459,0.0001277835],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","study_design_scores_codex":[0.0005599124,0.0001967686,0.02360231,0.0003493263,0.0004926367,0.00001534103,0.1846842,0.00009179163,0.002203462,0.3850585,0.3251846,0.07756113],"study_design_scores_gemma":[0.001453215,0.0003393249,0.0004126945,0.0003103018,0.00005431174,0.00002197071,0.01616817,0.00006692574,0.00007913939,0.002175134,0.9786981,0.0002206553],"study_design_candidate":"not_applicable","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9100438,0.004426115,0.001425578,0.001268849,0.004249532,0.0004057327,0.00007097136,0.000105465,0.078004],"genre_scores_gemma":[0.9826685,0.0006117992,0.002646427,0.0003440773,0.004263917,0.000004311433,0.000005586987,0.00004343552,0.00941193],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.6535136,"threshold_uncertainty_score":0.8840355,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.04396353817700301,"score_gpt":0.3179872963478167,"score_spread":0.2740237581708137,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}