{"id":"W6987819074","doi":"","title":"Utilizing NLP Sentiment Analysis Approach to Categorize Amazon Reviews against an Extended Testing Set","year":2024,"lang":"en","type":"article","venue":"Global Society of Scientific Research and Researchers - International Journal of Computer","topic":"Sentiment Analysis and Opinion Mining","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Calgary","funders":"","keywords":"Sentiment analysis; Random forest; Preprocessor; Categorization; Support vector machine; Bag-of-words model; Set (abstract data type); Feature (linguistics); Feature extraction; Product (mathematics)","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001966537,0.0009557834,0.0006190844,0.003101289,0.000399317,0.0009842297,0.0005186331,0.0004289019,0.002117752],"category_scores_gemma":[0.009780237,0.0001166412,0.0006963474,0.001780047,0.0002498716,0.0008171253,0.0006341572,0.0006009543,0.002117961],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0005862379,"about_ca_system_score_gemma":0.0005747728,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.004188361,"about_ca_topic_score_gemma":0.006549393,"domain_scores_codex":[0.9979582,0.0003972286,0.0003061149,0.0005065321,0.0007343818,0.00009757704],"domain_scores_gemma":[0.9936014,0.002521605,0.0007712627,0.0004565901,0.002522683,0.0001265453],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.001712377,0.001158371,0.2009509,0.001203638,0.0004922287,0.001334817,0.00153068,0.01094633,0.0696298,0.001800057,0.02841048,0.6808303],"study_design_scores_gemma":[0.0001456048,0.001810457,0.3493844,0.0002026653,0.0003173003,0.001448892,0.003758064,0.5371829,0.06946796,0.002780616,0.0333557,0.0001452977],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.8990248,0.0007364457,0.06656042,0.0005235252,0.0002906862,0.001566618,0.01664011,0.002231922,0.01242557],"genre_scores_gemma":[0.9032556,0.0001766664,0.07435103,0.0001693166,0.0001168419,0.0009929246,0.01757079,0.0001001591,0.003266693],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.004188361,"threshold_uncertainty_score":0.01040018,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.2396509484240729,"score_gpt":0.4432112019641853,"score_spread":0.2035602535401125,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}