{"id":"W4409576520","doi":"10.61091/jcmcc127a-145","title":"A Study of Modeling Algorithm Improvement for Semantic Analysis of Public English Texts","year":2025,"lang":"en","type":"article","venue":"Journal of Combinatorial Mathematics and Combinatorial Computing","topic":"Educational Systems and Policies","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Computer science; Natural language processing; Semantic analysis (machine learning); Artificial intelligence; Algorithm","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":true,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002556103,0.0002044197,0.001116146,0.0008308061,0.0001688724,0.0001761924,0.0007581855,0.00008760815,7.142687e-7],"category_scores_gemma":[0.0005363576,0.0001747667,0.0002991599,0.0013304,0.0000340613,0.0002278816,0.0002980622,0.0001692651,3.950075e-8],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.00006389813,"about_ca_system_score_gemma":0.0002838806,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00006031388,"about_ca_topic_score_gemma":0.000002167426,"domain_scores_codex":[0.9971105,0.0001054266,0.001724909,0.0002062587,0.0006094231,0.000243408],"domain_scores_gemma":[0.9949672,0.000924724,0.001495373,0.0003555086,0.002160022,0.00009714309],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0000124937,0.002476305,0.0007451744,0.0002942486,0.001966452,0.000001174008,0.01488542,0.001169619,0.0001751922,0.973281,0.00005265006,0.004940239],"study_design_scores_gemma":[0.005698814,0.00234064,0.0004696798,0.0004254778,0.001325394,0.000004246873,0.006266379,0.6318513,0.0004336209,0.3507797,0.00009939186,0.0003054326],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.7180763,0.0001987186,0.2739227,0.00006111609,0.007233275,0.0003867092,0.00000256641,0.00001205886,0.000106561],"genre_scores_gemma":[0.9893248,0.000008928401,0.01025662,0.000008486916,0.0003808116,0.00000588751,8.026333e-7,0.000009068399,0.000004592249],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.6306816,"threshold_uncertainty_score":0.712678,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.02292818587694894,"score_gpt":0.2943475408140213,"score_spread":0.2714193549370723,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}