{"id":"W7106786599","doi":"10.48448/wcv1-dj44","title":"Improving Context Fidelity via Native Retrieval-Augmented Reasoning","year":2025,"lang":"","type":"other","venue":"Open MIND","topic":"","field":"","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"Mila - Quebec Artificial Intelligence Institute; McGill University","funders":"","keywords":"Counterfactual thinking; Context (archaeology); Process (computing); Fidelity; Case-based reasoning; Scientific reasoning","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004506174,0.001444292,0.001114195,0.001141181,0.0007473172,0.00250864,0.003577437,0.001879021,0.004436119],"category_scores_gemma":[0.02377018,0.0007652197,0.00130862,0.0006868778,0.001464775,0.004393556,0.004888147,0.002755512,0.001697465],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001042261,"about_ca_system_score_gemma":0.002015079,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.003971194,"about_ca_topic_score_gemma":0.007199863,"domain_scores_codex":[0.9953056,0.00230932,0.0002736227,0.001068836,0.000834003,0.0002086115],"domain_scores_gemma":[0.990903,0.004464322,0.0005670933,0.003082479,0.000823438,0.0001596594],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0005581902,0.0006792529,0.004086318,0.0006596246,0.0002263065,0.000544972,0.0009915303,0.2591953,0.02968363,0.03516043,0.01353532,0.6546791],"study_design_scores_gemma":[0.0001078596,0.0001152346,0.000402002,0.00005125448,0.00007264465,0.0001773996,0.0001068602,0.9419558,0.01186764,0.0384139,0.006684507,0.00004489088],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.03003518,0.0008247785,0.9494615,0.0007295034,0.00009831878,0.0002276957,0.0003304468,0.01369906,0.004593498],"genre_scores_gemma":[0.524173,0.0002369665,0.4696138,0.0008003319,0.00009431559,0.0002167826,0.0008960617,0.0008114257,0.003157281],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.004506174,"threshold_uncertainty_score":0.02383119,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.02829326156452239,"score_gpt":0.3227272862862754,"score_spread":0.294434024721753,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}