{"id":"W4400484677","doi":"10.1145/3663529.3663857","title":"Decoding Anomalies! Unraveling Operational Challenges in Human-in-the-Loop Anomaly Validation","year":2024,"lang":"en","type":"article","venue":"","topic":"Software System Performance and Reliability","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"Concordia University","funders":"","keywords":"Decoding methods; Computer science; Anomaly (physics); Loop (graph theory); Anomaly detection; Data mining; Algorithm; Mathematics; Physics","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04518735,0.001413502,0.001102981,0.002129105,0.001554867,0.006676998,0.003449351,0.002429853,0.001643836],"category_scores_gemma":[0.2045261,0.0006918244,0.0007092117,0.001251841,0.005144448,0.01077291,0.005585347,0.005107566,0.0008655549],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.002090435,"about_ca_system_score_gemma":0.005427221,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00574676,"about_ca_topic_score_gemma":0.005983636,"domain_scores_codex":[0.9421424,0.03669704,0.002601206,0.004123076,0.01284172,0.001594474],"domain_scores_gemma":[0.780953,0.1461096,0.01186132,0.03352818,0.0246933,0.002854748],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"qualitative","study_design_scores_codex":[0.000828728,0.001082517,0.05894807,0.001288321,0.0002926021,0.001149583,0.02183168,0.1241698,0.04592667,0.09623241,0.01395847,0.6342912],"study_design_scores_gemma":[0.00009799257,0.0006284052,0.009253163,0.0006248191,0.00007603235,0.0006821677,0.006717307,0.7131479,0.03619692,0.2017659,0.03055293,0.0002564734],"study_design_candidate":"qualitative","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.0827215,0.000468192,0.8961267,0.01070299,0.0001448327,0.0003264268,0.0002468661,0.005002601,0.004259878],"genre_scores_gemma":[0.6213383,0.0002233544,0.3751537,0.001176468,0.0000604822,0.0002167222,0.0004915951,0.0006180502,0.0007212813],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.04518735,"threshold_uncertainty_score":0.2389764,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.06021226248420644,"score_gpt":0.3055387277811504,"score_spread":0.245326465296944,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}