{"id":"W2604394978","doi":"10.18853/jjell.2008.50.3.007","title":"Reconsideration of the Development of Rating Criteria in Operational Language Testing Contexts","year":2008,"lang":"en","type":"article","venue":"The Jungang Journal of English Language and Literature","topic":"EFL/ESL Teaching and Learning","field":"Arts and Humanities","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Psychology; Language assessment; Linguistics; Computer science; Mathematics education; Philosophy","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.2962428,0.001171949,0.001862124,0.005503049,0.00362478,0.01058467,0.006927748,0.002624168,0.001070859],"category_scores_gemma":[0.5261641,0.001197702,0.001396049,0.004520944,0.02070755,0.01227253,0.006355709,0.007249326,0.0004716186],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.005988159,"about_ca_system_score_gemma":0.01194392,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.009370795,"about_ca_topic_score_gemma":0.007987844,"domain_scores_codex":[0.6578676,0.2553683,0.03437185,0.007893356,0.0412837,0.003215241],"domain_scores_gemma":[0.3345381,0.5091822,0.02244642,0.02961211,0.09905697,0.005164199],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.0002753403,0.0004782843,0.05707985,0.002423099,0.0002194254,0.000593843,0.09642006,0.003431491,0.004348188,0.4757425,0.008927797,0.3500601],"study_design_scores_gemma":[0.0003370001,0.001633649,0.07903237,0.00904836,0.0003185481,0.002188349,0.1057383,0.04617175,0.007852331,0.5879406,0.1587115,0.001027139],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.1404725,0.005355904,0.7634038,0.03169798,0.002350457,0.001588839,0.0001988765,0.0004769986,0.05445472],"genre_scores_gemma":[0.6761166,0.000821308,0.3182214,0.001877057,0.0003961065,0.001227323,0.0001068497,0.0002013322,0.00103211],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.2962428,"threshold_uncertainty_score":0.8678579,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.0236169268155502,"score_gpt":0.2502907647183861,"score_spread":0.2266738379028359,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}