{"id":"W6948291587","doi":"10.48448/1c51-n793","title":"What if...?: Thinking Counterfactual Keywords Helps to Mitigate Hallucination in Large Multi-modal Models","year":2024,"lang":"en","type":"other","venue":"Underline Science Inc.","topic":"Nephrotoxicity and Medicinal Plants","field":"Medicine","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Counterfactual thinking; Counterfactual conditional; Context (archaeology); Cognition; Bridging (networking); Constraint (computer-aided design)","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002034729,0.0006957375,0.0003139497,0.000314827,0.0004064667,0.001677897,0.0008946574,0.0008765537,0.004436841],"category_scores_gemma":[0.01613687,0.0002162016,0.0007149997,0.0001742238,0.001390456,0.003255784,0.001903827,0.001016425,0.0005443562],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0003729643,"about_ca_system_score_gemma":0.000811526,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001108478,"about_ca_topic_score_gemma":0.001461098,"domain_scores_codex":[0.999008,0.0004757684,0.00005632172,0.0002166181,0.000183549,0.00005982371],"domain_scores_gemma":[0.9930995,0.004221199,0.0006631717,0.001386937,0.0004148222,0.0002143653],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.003063329,0.0004288814,0.02403286,0.001721013,0.0003846426,0.001873563,0.009661498,0.225205,0.1514676,0.2218036,0.006068601,0.3542895],"study_design_scores_gemma":[0.0000892219,0.0003221872,0.002731206,0.0001216726,0.0001257826,0.0005475417,0.001239914,0.756673,0.03719179,0.191692,0.00915514,0.0001105759],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.1872658,0.000171803,0.7998062,0.001324544,0.00007921052,0.0001282352,0.0002836646,0.001879792,0.009060774],"genre_scores_gemma":[0.8612394,0.0000820649,0.1366554,0.0002198735,0.00001885371,0.00007824222,0.000192921,0.0001891348,0.001324137],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.004436841,"threshold_uncertainty_score":0.01484275,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.03869328472171494,"score_gpt":0.3418334375089332,"score_spread":0.3031401527872183,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}