{"id":"W7108248275","doi":"","title":"CRAwDAD: Causal Reasoning Augmentation with Dual-Agent Debate","year":2025,"lang":"","type":"article","venue":"arXiv (Cornell University)","topic":"Explainable Artificial Intelligence (XAI)","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"Alberta Children's Hospital; University of Calgary","funders":"","keywords":"Counterfactual thinking; Counterfactual conditional; Causal reasoning; Causal inference; Causal model; Inference; Natural language understanding; Philosophy of language; Natural language; Causation","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaepi_narrow"],"consensus_categories":[],"category_scores_codex":[0.0005679622,0.0005182486,0.0004205485,0.0006084715,0.0010287,0.0005155751,0.001232424,0.0002135714,0.0002134283],"category_scores_gemma":[0.0000951305,0.0005931808,0.0001653693,0.003920998,0.0004119248,0.002047,0.0007297283,0.000461028,0.0005557029],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.000807556,"about_ca_system_score_gemma":0.0007092791,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001144346,"about_ca_topic_score_gemma":0.0008720144,"domain_scores_codex":[0.9962243,0.0003641604,0.0004607413,0.001702433,0.0002372784,0.00101107],"domain_scores_gemma":[0.9972808,0.0002365513,0.0003395317,0.001308471,0.0005028427,0.0003317997],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0003869615,0.0003701811,0.009708243,0.0001006449,0.0003088432,0.002782212,0.002121202,0.2307914,0.0007539652,0.7455817,0.0007887009,0.006305985],"study_design_scores_gemma":[0.001105231,0.0005641479,0.002937552,0.0005856856,0.0003327458,0.000031111,0.002935766,0.9544689,0.01892527,0.0148522,0.002282624,0.0009787842],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.3541699,0.0001103873,0.6349204,0.0003916575,0.00076795,0.0004224699,0.000004432283,0.0001751578,0.009037701],"genre_scores_gemma":[0.9779457,0.000231167,0.001752903,0.000443302,0.00007003913,0.000002253842,0.000006859559,0.00002643852,0.01952127],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.7307295,"threshold_uncertainty_score":0.999652,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.05426226973516782,"score_gpt":0.2104052852759768,"score_spread":0.156143015540809,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}