{"id":"W33598619","doi":"","title":"Improving a statistical language model by modulating the effects of context words","year":2008,"lang":"en","type":"article","venue":"UCL Discovery (University College London)","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Feature (linguistics); Word (group theory); Artificial intelligence; Context (archaeology); Feature vector; Language model; Artificial neural network; Context model; Natural language processing; Pattern recognition (psychology); Mathematics; Linguistics","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0001248324,0.0001504508,0.0002328313,0.0000845068,0.0003637392,0.00003123561,0.0008571977,0.00006264683,0.000004801997],"category_scores_gemma":[0.0000976059,0.0001371964,0.00008902671,0.0003306993,0.0001350333,0.001068296,0.0005116574,0.0001881006,0.000005084049],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0001002266,"about_ca_system_score_gemma":0.0001571052,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.0006827853,"about_ca_topic_score_gemma":0.00005906576,"domain_scores_codex":[0.9987466,0.0001041078,0.0001519143,0.0003704969,0.0003342756,0.0002925504],"domain_scores_gemma":[0.9987851,0.0004103515,0.0001241883,0.0005568092,0.00004487269,0.00007866426],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.00024416,0.0005327219,0.002718626,0.000530859,0.000235163,0.001706844,0.03431768,0.02294449,0.03716224,0.8382518,0.01100128,0.05035417],"study_design_scores_gemma":[0.000821248,0.00005401321,0.000371589,0.00002999681,0.00001712761,0.00001678055,0.001028388,0.996555,0.0006257306,0.0002612465,0.00004881925,0.0001701126],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.4320532,0.0001148831,0.5670999,0.000115566,0.00007981377,0.0001574737,0.00007845868,0.00004230392,0.0002583868],"genre_scores_gemma":[0.9747686,0.0000102896,0.02073132,0.0001482111,0.00001825947,0.000001021516,0.000004716338,0.000009219103,0.004308308],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.9736105,"threshold_uncertainty_score":0.5594707,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.006199984077451072,"score_gpt":0.1854097546974772,"score_spread":0.1792097706200262,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}