{"id":"W1979004204","doi":"10.1145/1557626.1557630","title":"Issues in pattern mining and their resolutions","year":2009,"lang":"en","type":"article","venue":"","topic":"Data Mining Algorithms and Applications","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"ca_institutions":"Concordia University","funders":"","keywords":"Overfitting; Computer science; Data mining; Reliability (semiconductor); Anomaly detection; Sequential Pattern Mining; Anomaly (physics); Data science; Machine learning; Artificial neural network","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02200253,0.001502726,0.002520986,0.005688096,0.003523886,0.01362744,0.004808167,0.00637722,0.005230519],"category_scores_gemma":[0.114259,0.00213316,0.002534449,0.01153615,0.00935232,0.0304822,0.007105845,0.01180791,0.002403937],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001942566,"about_ca_system_score_gemma":0.001769433,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00108291,"about_ca_topic_score_gemma":0.0004541202,"domain_scores_codex":[0.9716627,0.01132664,0.002992196,0.004352292,0.00869529,0.0009709377],"domain_scores_gemma":[0.9279494,0.05480684,0.0044712,0.00707407,0.004739917,0.0009585955],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.00007237873,0.00008493674,0.001848458,0.0006514636,0.0001125907,0.0006100396,0.0007051391,0.004266202,0.0003248266,0.8123091,0.01347008,0.1655447],"study_design_scores_gemma":[0.0000124501,0.00001492274,0.0002952924,0.0001560932,0.00001638783,0.0005255393,0.0003325931,0.008917524,0.0003043753,0.9672527,0.02214749,0.00002468483],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.01144431,0.07133307,0.7949156,0.09134134,0.003237999,0.0002973926,0.0007556944,0.0004936648,0.02618095],"genre_scores_gemma":[0.2711764,0.06964529,0.617174,0.00824407,0.01460205,0.001040865,0.001582351,0.0002489484,0.01628603],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.02200253,"threshold_uncertainty_score":0.1163619,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.0181417059000252,"score_gpt":0.2675401528924545,"score_spread":0.2493984469924293,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}