{"id":"W4409795195","doi":"10.61091/jcmcc127b-515","title":"An Empirical Study of Intelligent Algorithms for Evaluating English Teaching Effectiveness in Colleges and Universities","year":2025,"lang":"en","type":"article","venue":"Journal of Combinatorial Mathematics and Combinatorial Computing","topic":"Educational Technology and Pedagogy","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Computer science; Empirical research; Mathematics education; College English; Artificial intelligence; Engineering management; Engineering; Psychology; Mathematics; Statistics","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":true,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01001817,0.0005692049,0.0005622799,0.003059249,0.0004694343,0.002062488,0.0007125848,0.0007261552,0.0007397925],"category_scores_gemma":[0.06511784,0.000187702,0.0004749326,0.003132358,0.0007232694,0.00248051,0.000562862,0.0005777968,0.0001535331],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001156057,"about_ca_system_score_gemma":0.001286445,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.002801696,"about_ca_topic_score_gemma":0.001873065,"domain_scores_codex":[0.9917867,0.004200676,0.0008882002,0.0005935654,0.002284564,0.0002463273],"domain_scores_gemma":[0.9601524,0.02875996,0.00334006,0.001741819,0.005430974,0.0005747535],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"observational","study_design_scores_codex":[0.0003894071,0.0008503144,0.427197,0.0007503834,0.0004680491,0.000205251,0.001969479,0.09003278,0.001993668,0.01321981,0.002009345,0.4609145],"study_design_scores_gemma":[0.00007113968,0.001071764,0.2363587,0.0003014026,0.0002609555,0.0004099912,0.002137878,0.7422968,0.005633516,0.007208803,0.004163521,0.00008566079],"study_design_candidate":"observational","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.8613266,0.002129791,0.1227792,0.0004234338,0.00006847468,0.0003850783,0.0002454248,0.0002381573,0.01240379],"genre_scores_gemma":[0.9620708,0.0003695773,0.03675292,0.00003257787,0.00001600753,0.0001285675,0.0001532019,0.00001709481,0.0004590951],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.01001817,"threshold_uncertainty_score":0.05298179,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.03690614867721735,"score_gpt":0.3878501029996974,"score_spread":0.35094395432248,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}