{"id":"W7153004372","doi":"10.63282/3117-5481/aijcst-v6i6p108","title":"Integrating Observability, Defect Prediction, and Decision Intelligence for Reliable AI-Driven Software Systems","year":2024,"lang":"","type":"article","venue":"American International Journal of Computer Science and Technology","topic":"Software Reliability and Analysis Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Software quality; Software system; Software quality control; Software; Software quality analyst; Software development; Software architecture; Verification and validation; Debugging","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005231268,0.0008246291,0.0005624857,0.002877362,0.0008890214,0.005620884,0.001414588,0.00127052,0.00101298],"category_scores_gemma":[0.01515165,0.0006618993,0.0007746174,0.001382498,0.004213293,0.006935196,0.003597605,0.002710043,0.000270384],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.002584309,"about_ca_system_score_gemma":0.003383893,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.004442445,"about_ca_topic_score_gemma":0.003799609,"domain_scores_codex":[0.9961842,0.001186147,0.000369292,0.000568996,0.001399506,0.0002918104],"domain_scores_gemma":[0.9873312,0.00661149,0.002319961,0.001682322,0.001503605,0.0005513569],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.000089936,0.000190283,0.01556313,0.0003692568,0.0001448481,0.0003272341,0.002296015,0.2427438,0.004393294,0.5463215,0.001213156,0.1863476],"study_design_scores_gemma":[0.00001270552,0.0001312444,0.002921462,0.0001851134,0.0000704664,0.0001078427,0.0006991966,0.5621433,0.001927805,0.4255129,0.006214553,0.00007345874],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.04124722,0.0008148844,0.9455915,0.003282397,0.00004223247,0.00009321057,0.00004468992,0.0004671609,0.008416824],"genre_scores_gemma":[0.8248789,0.0008036358,0.1729078,0.0001892486,0.00006448801,0.00009123994,0.00008834073,0.00004860296,0.000927905],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.005620884,"threshold_uncertainty_score":0.02766591,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.01569497516323298,"score_gpt":0.321151208990072,"score_spread":0.305456233826839,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}