{"id":"W2490032828","doi":"10.4018/978-1-60566-766-9.ch015","title":"Machine Learning Applications in Mega-Text Processing","year":2010,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Sentiment Analysis and Opinion Mining","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Ottawa; Agricultural Research Institute of Ontario","funders":"","keywords":"Computer science; Focus (optics); Selection (genetic algorithm); Feature selection; Text processing; The Internet; Prime (order theory); Feature (linguistics); Artificial intelligence; Representation (politics); Natural language processing; Word (group theory); Sentiment analysis; Data science; Information retrieval; World Wide Web; Linguistics","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaepi_narrow"],"consensus_categories":[],"category_scores_codex":[0.0001907624,0.0003184509,0.0003816391,0.0001766695,0.0002090946,0.0003433666,0.0008670456,0.0002866926,0.00004363387],"category_scores_gemma":[0.000009379753,0.0003077536,0.0001704719,0.00007993144,0.00005438823,0.0001180203,0.0003326792,0.0006442827,0.0001118176],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.00009394083,"about_ca_system_score_gemma":0.0001339491,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00004794783,"about_ca_topic_score_gemma":0.0001138169,"domain_scores_codex":[0.9982658,0.00001672478,0.0004180407,0.0006350033,0.0003764671,0.0002879844],"domain_scores_gemma":[0.99896,0.00003032222,0.0003159585,0.0004936011,0.00009116979,0.00010898],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","study_design_scores_codex":[0.00000149522,0.000008697506,0.000287003,0.0000147294,0.00002174177,0.000008814218,0.0001511331,0.00002177958,0.0000361605,0.9454007,0.00002723393,0.05402058],"study_design_scores_gemma":[0.0006861697,0.00007088573,0.0002222331,0.000439022,0.0001050655,0.00003633615,0.00003927724,0.03212199,0.000213972,0.3737914,0.5908247,0.001448945],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"genre_codex":"other","genre_gemma":"empirical","genre_scores_codex":[0.00004556012,0.000946838,0.048601,0.00007810774,0.0001351571,0.0002273621,0.000004679993,0.0001678042,0.9497935],"genre_scores_gemma":[0.5635232,0.00005016307,0.06976242,0.0008401472,0.0009986737,0.0001490317,0.00007756624,0.0001461731,0.3644526],"genre_candidate":"other","genre_consensus":null,"teacher_disagreement_score":0.5907975,"threshold_uncertainty_score":0.9999375,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.01676752531484623,"score_gpt":0.2603800503594275,"score_spread":0.2436125250445813,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}