{"id":"W4409795195","doi":"10.61091/jcmcc127b-515","title":"An Empirical Study of Intelligent Algorithms for Evaluating English Teaching Effectiveness in Colleges and Universities","year":2025,"lang":"en","type":"article","venue":"Journal of Combinatorial Mathematics and Combinatorial Computing","topic":"Educational Technology and Pedagogy","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Computer science; Empirical research; Mathematics education; College English; Artificial intelligence; Engineering management; Engineering; Psychology; Mathematics; Statistics","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":true,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003943239,0.0001576888,0.0005694559,0.0004364011,0.0002379774,0.0000977716,0.0005137366,0.0001263517,2.834578e-7],"category_scores_gemma":[0.0006698279,0.0001500696,0.0000564318,0.0002955866,0.00006795823,0.0002685972,0.0002428882,0.0003901266,2.309212e-8],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.00009834478,"about_ca_system_score_gemma":0.0003253615,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00001689387,"about_ca_topic_score_gemma":0.000001111865,"domain_scores_codex":[0.9982369,0.0003836926,0.0007111392,0.0002168122,0.0002778664,0.0001736063],"domain_scores_gemma":[0.9956307,0.003010916,0.0004953813,0.0001939495,0.0006103313,0.00005876436],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.0001237863,0.00210344,0.007653011,0.0002697621,0.0001063734,0.000004825248,0.01902269,0.000164492,0.0002040847,0.965121,0.00001028421,0.005216299],"study_design_scores_gemma":[0.008107638,0.005953318,0.006100297,0.0007589139,0.0001076167,0.00002027394,0.03312628,0.07041708,0.001059784,0.8740147,0.00005992596,0.0002741646],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9479703,0.0001592728,0.04618927,0.00006344656,0.00512804,0.0003967328,7.154417e-7,0.00001947164,0.00007274031],"genre_scores_gemma":[0.9847685,0.000006393667,0.0150388,0.000007623906,0.000165689,0.000003866824,2.969147e-7,0.000006640265,0.00000216315],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.09110624,"threshold_uncertainty_score":0.6119658,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.03690614867721735,"score_gpt":0.3878501029996974,"score_spread":0.35094395432248,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}