{"id":"W4412703901","doi":"10.1145/3696630.3728522","title":"Can Hessian-Based Insights Support Fault Diagnosis in Attention-based Models?","year":2025,"lang":"en","type":"article","venue":"","topic":"Explainable Artificial Intelligence (XAI)","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"Dalhousie University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Hessian matrix; Computer science; Fault (geology); Artificial intelligence; Machine learning; Geology; Mathematics; Applied mathematics","routes":{"ca_aff":true,"ca_fund":true,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001253915,0.00140609,0.0007126365,0.001187134,0.0003197473,0.001173446,0.00103159,0.001019605,0.002476109],"category_scores_gemma":[0.01261154,0.0004481259,0.0004126461,0.0004146498,0.0007521671,0.002679293,0.001073329,0.001423495,0.0003521028],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001087422,"about_ca_system_score_gemma":0.001060626,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.009111614,"about_ca_topic_score_gemma":0.01328764,"domain_scores_codex":[0.9996423,0.00008701198,0.00002330069,0.00009454146,0.00009509939,0.00005770771],"domain_scores_gemma":[0.9956567,0.002446686,0.0006650411,0.0004107172,0.0006261566,0.0001947565],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0003370808,0.0001581461,0.02219054,0.0002062038,0.0001170692,0.0002599894,0.0002710414,0.8208982,0.01634342,0.01068528,0.002589852,0.1259432],"study_design_scores_gemma":[0.000004500713,0.00002730507,0.0009563642,0.00000642097,0.000007770545,0.00001961649,0.00001463405,0.9919577,0.001511213,0.005386323,0.0001019635,0.000006214509],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.3665065,0.0006801457,0.6252948,0.001427606,0.00007551533,0.00005346761,0.0003650274,0.002742321,0.002854573],"genre_scores_gemma":[0.9743601,0.00006703202,0.0247861,0.00007811892,0.00001269356,0.00001142072,0.0001270519,0.00006740021,0.0004901098],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.009111614,"threshold_uncertainty_score":0.01811713,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.02732413311328661,"score_gpt":0.2758546596971958,"score_spread":0.2485305265839092,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}