{"id":"W1253067","doi":"","title":"Spam Detection: A Syntax and Semantic-Based Approach","year":2006,"lang":"en","type":"article","venue":"IKE","topic":"Spam and Phishing Detection","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Computer science; Syntax; Natural language processing; Artificial intelligence; Semantics (computer science); Linguistics; Programming language; Philosophy","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004978177,0.001262402,0.001520757,0.01391777,0.001795059,0.004082025,0.001863842,0.002009805,0.002623941],"category_scores_gemma":[0.009549716,0.0005663937,0.001818861,0.005538113,0.001883497,0.006623514,0.00284358,0.001088619,0.001581754],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001265714,"about_ca_system_score_gemma":0.001868002,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.002775164,"about_ca_topic_score_gemma":0.003166083,"domain_scores_codex":[0.9954988,0.002003418,0.0004685904,0.0008562969,0.0009564295,0.0002163873],"domain_scores_gemma":[0.993434,0.002726791,0.0006501434,0.0006818937,0.002311426,0.0001957423],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.001244809,0.001135151,0.05390375,0.001696498,0.0005459042,0.001132893,0.003133464,0.01298864,0.02784692,0.0682105,0.01582259,0.8123388],"study_design_scores_gemma":[0.0002219728,0.000759821,0.0370836,0.0003836301,0.0008939391,0.003214765,0.007230921,0.687916,0.02217662,0.2060724,0.03365419,0.000392069],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.08534477,0.000674483,0.8929341,0.001539791,0.0001982009,0.0008563725,0.003657192,0.006929659,0.007865504],"genre_scores_gemma":[0.3937832,0.0003452619,0.5967898,0.0004319627,0.0001893314,0.0005197388,0.005260702,0.0003420292,0.002337938],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.01391777,"threshold_uncertainty_score":0.02632743,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.006543711778088978,"score_gpt":0.1806008433942105,"score_spread":0.1740571316161215,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}