{"id":"W4395021982","doi":"10.48550/arxiv.2404.12535","title":"Is There No Such Thing as a Bad Question? H4R: HalluciBot For Ratiocination, Rewriting, Ranking, and Routing","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Psychedelics and Drug Studies","field":"Psychology","cited_by":2,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"Institute for Catastrophic Loss Reduction","keywords":"Epistemology; Psychology; Philosophy","routes":{"ca_aff":false,"ca_fund":true,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01120328,0.0009098983,0.001014875,0.0004223945,0.001076367,0.003174667,0.001714532,0.002053555,0.008976287],"category_scores_gemma":[0.07033122,0.0004871979,0.001188717,0.0004429302,0.002352604,0.005910246,0.002205259,0.002045265,0.003128363],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001719059,"about_ca_system_score_gemma":0.002433536,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.005303768,"about_ca_topic_score_gemma":0.003756436,"domain_scores_codex":[0.9915918,0.004381264,0.0003613059,0.001815824,0.001225357,0.0006243858],"domain_scores_gemma":[0.9616053,0.02570867,0.002925526,0.005331506,0.003216748,0.001212297],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.004049258,0.001330199,0.1786371,0.001849472,0.0005791577,0.001986941,0.01335748,0.1316743,0.06295037,0.165634,0.05990789,0.3780439],"study_design_scores_gemma":[0.0003548161,0.001448286,0.02203969,0.0001569158,0.0001785556,0.00136401,0.003460148,0.7877518,0.03059526,0.1243452,0.02807011,0.0002352344],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.3994056,0.0004917326,0.5402009,0.01006493,0.0003337992,0.00131018,0.002559305,0.007867364,0.03776614],"genre_scores_gemma":[0.9286659,0.00007982043,0.06180478,0.001764926,0.00006848825,0.000444634,0.001378618,0.0004359895,0.005356805],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.01120328,"threshold_uncertainty_score":0.05924934,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.07943446247142871,"score_gpt":0.2821092111644805,"score_spread":0.2026747486930518,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}