{"id":"W4412636726","doi":"10.1007/978-3-031-99264-3_36","title":"Making Generative AI Hallucinations Useful by Reassessing the Troublemaker Agent Strategy","year":2025,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Computability, Logic, AI Algorithms","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"ca_institutions":"McGill University","funders":"","keywords":"Generative grammar; Computer science; Artificial intelligence","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001388388,0.0004995213,0.0002545063,0.0004192856,0.001076534,0.00426258,0.001338772,0.00195128,0.008327536],"category_scores_gemma":[0.01204842,0.0002261088,0.0002549415,0.0002641857,0.008429569,0.008858246,0.002570382,0.004092562,0.002046289],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0006810637,"about_ca_system_score_gemma":0.0004556415,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.000451356,"about_ca_topic_score_gemma":0.0005407337,"domain_scores_codex":[0.99915,0.0004969623,0.00002083661,0.00009489198,0.0001733148,0.00006392192],"domain_scores_gemma":[0.996438,0.002254555,0.0001734785,0.0006616405,0.0002594471,0.0002129136],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0000976982,0.00003322915,0.0004144833,0.0000918796,0.00001569004,0.0003321944,0.009437416,0.001935233,0.002690677,0.9127674,0.01670967,0.05547443],"study_design_scores_gemma":[0.00003129019,0.00002835535,0.0001505889,0.0000429119,0.00001129265,0.0003790183,0.001751033,0.008727087,0.001498823,0.9369552,0.05040071,0.00002362239],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.0648998,0.00299849,0.4772344,0.05090856,0.002813371,0.00009843117,0.0001276114,0.001981664,0.3989377],"genre_scores_gemma":[0.9049338,0.0008216547,0.054839,0.003966456,0.0005518015,0.0000633277,0.00007100942,0.0005584505,0.03419444],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.008327536,"threshold_uncertainty_score":0.02785844,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.08497441642269019,"score_gpt":0.349361195893825,"score_spread":0.2643867794711349,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}