{"id":"W7113200082","doi":"","title":"Hallucinations in Large Foundation Models: Characterization, Quantification, Detection, Avoidance, and Mitigation","year":2025,"lang":"en","type":"article","venue":"Scholar Commons (University of South Carolina)","topic":"Explainable Artificial Intelligence (XAI)","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Deception; Generative grammar; Foundation (evidence); Unintended consequences; Scope (computer science); Argument (complex analysis); Perception; Salient","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006925628,0.0001069855,0.0001625792,0.0006601147,0.0007032156,0.0001031627,0.0004814679,0.000102878,0.00000633583],"category_scores_gemma":[0.0001217354,0.0001526552,0.00004330019,0.00152849,0.00009621493,0.002332721,0.0001856789,0.0002481337,0.00001961213],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0001040157,"about_ca_system_score_gemma":0.0001089846,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.0001337548,"about_ca_topic_score_gemma":0.001480195,"domain_scores_codex":[0.9988815,0.0001777499,0.000207661,0.0003621845,0.0001813276,0.0001896148],"domain_scores_gemma":[0.9987642,0.00005926345,0.0001731792,0.0004548968,0.0004957775,0.00005265254],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0000461992,0.0002675231,0.05664393,0.000121998,0.00004890852,0.000008005663,0.02171851,0.00261095,0.02301049,0.8659409,0.00003782128,0.02954479],"study_design_scores_gemma":[0.0007160686,0.00004760092,0.3672117,0.0001310874,0.00003199717,0.000003151172,0.005311508,0.5788528,0.005915027,0.03823606,0.003223417,0.0003196435],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.3884168,0.00003465393,0.6094587,0.001259312,0.00008836179,0.0001697868,0.00001033986,0.00005793475,0.0005041144],"genre_scores_gemma":[0.9949136,0.00003915371,0.004623062,0.00005055971,0.00000760144,0.000002270321,0.00003847265,0.000004927637,0.0003203286],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.8277048,"threshold_uncertainty_score":0.6225097,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.01595243243373195,"score_gpt":0.2250382083375209,"score_spread":0.2090857759037889,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}