{"id":"W2912504907","doi":"10.22190/jtesap1803525s","title":"IS ESP STANDARDIZED ASSESSMENT FEASIBLE?","year":2019,"lang":"en","type":"article","venue":"Journal of Teaching English for Specific and Academic Purposes","topic":"EFL/ESL Teaching and Learning","field":"Arts and Humanities","cited_by":2,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Benchmarking; Standardized test; Quality assessment; Best practice; Quality assurance; Educational assessment; Language assessment; Variety (cybernetics); Quality (philosophy); Psychology; Computer science; Evaluation methods; Political science; Business; Pedagogy; External quality assessment; Engineering; Mathematics education; Operations management; Artificial intelligence; Marketing","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.1251345,0.0006723074,0.00114457,0.00374242,0.002336574,0.01904657,0.003782722,0.00468212,0.004998062],"category_scores_gemma":[0.2616304,0.0006381852,0.001218743,0.007437893,0.01373324,0.02596683,0.01160539,0.004759015,0.002658078],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.006744479,"about_ca_system_score_gemma":0.02138278,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00582166,"about_ca_topic_score_gemma":0.003873635,"domain_scores_codex":[0.7713923,0.1433293,0.02259722,0.01034874,0.04721215,0.005120337],"domain_scores_gemma":[0.7838502,0.1087574,0.01760266,0.03161148,0.05337828,0.004799981],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","study_design_scores_codex":[0.0001011863,0.0001232183,0.01400948,0.001864361,0.00007134854,0.0002987453,0.01094629,0.002048899,0.0004106968,0.558122,0.01727683,0.394727],"study_design_scores_gemma":[0.00005140002,0.000319717,0.01252396,0.009239192,0.00004435136,0.0008453867,0.02884518,0.005029525,0.0008304665,0.6284458,0.3136916,0.0001334306],"study_design_candidate":"observational","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.0747617,0.02518592,0.4139164,0.1992719,0.004599621,0.0006419552,0.0006656213,0.001326458,0.2796304],"genre_scores_gemma":[0.8043081,0.01427825,0.1557604,0.01263292,0.001718042,0.001349648,0.001256643,0.000395803,0.00830016],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.1251345,"threshold_uncertainty_score":0.6617825,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.04396353817700301,"score_gpt":0.3179872963478167,"score_spread":0.2740237581708137,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}