{"id":"W7108248275","doi":"","title":"CRAwDAD: Causal Reasoning Augmentation with Dual-Agent Debate","year":2025,"lang":"","type":"article","venue":"arXiv (Cornell University)","topic":"Explainable Artificial Intelligence (XAI)","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"Alberta Children's Hospital; University of Calgary","funders":"","keywords":"Counterfactual thinking; Counterfactual conditional; Causal reasoning; Causal inference; Causal model; Inference; Natural language understanding; Philosophy of language; Natural language; Causation","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009396739,0.002239188,0.001483005,0.002220321,0.001780799,0.004202745,0.006267766,0.003893197,0.01437059],"category_scores_gemma":[0.03949907,0.0009196509,0.003450109,0.001342392,0.002510039,0.01036496,0.00873981,0.008318972,0.003563432],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.003191797,"about_ca_system_score_gemma":0.004961165,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.007613701,"about_ca_topic_score_gemma":0.01730545,"domain_scores_codex":[0.9913995,0.004465569,0.000456832,0.001714546,0.001664041,0.0002995972],"domain_scores_gemma":[0.9808494,0.01410154,0.0005826993,0.002867224,0.001101723,0.0004974465],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.001606122,0.001323462,0.006974854,0.002481134,0.0006241649,0.001138947,0.001859645,0.1626627,0.005986064,0.1660095,0.1513204,0.4980131],"study_design_scores_gemma":[0.0005242019,0.0001044779,0.0003923463,0.000166083,0.000108962,0.0001524845,0.0002213704,0.7630481,0.004665504,0.1774675,0.05308168,0.00006735943],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.03078681,0.00412179,0.8634911,0.01069162,0.001076642,0.001392735,0.009695245,0.05779032,0.02095368],"genre_scores_gemma":[0.2890059,0.0007846793,0.6798114,0.003659192,0.0003178369,0.001192219,0.01492591,0.001516676,0.008786206],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.01437059,"threshold_uncertainty_score":0.04969531,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.05426226973516782,"score_gpt":0.2104052852759768,"score_spread":0.156143015540809,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}