{"id":"W4395021982","doi":"10.48550/arxiv.2404.12535","title":"Is There No Such Thing as a Bad Question? H4R: HalluciBot For Ratiocination, Rewriting, Ranking, and Routing","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Psychedelics and Drug Studies","field":"Psychology","cited_by":2,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"Institute for Catastrophic Loss Reduction","keywords":"Epistemology; Psychology; Philosophy","routes":{"ca_aff":false,"ca_fund":true,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaepi_narrow"],"consensus_categories":[],"category_scores_codex":[0.0007568628,0.0003836183,0.0004097011,0.0002218342,0.0004292126,0.0001766838,0.0003527457,0.0004197517,0.0002339423],"category_scores_gemma":[0.0001606196,0.0004294216,0.0002349943,0.0002304149,0.00011444,0.00009140741,0.000694813,0.0006666228,0.000220323],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0001347394,"about_ca_system_score_gemma":0.00009683811,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.0009648876,"about_ca_topic_score_gemma":0.000184325,"domain_scores_codex":[0.9977629,0.0001478193,0.0003166235,0.001285802,0.00009063778,0.0003962077],"domain_scores_gemma":[0.998381,0.0002746275,0.0002856409,0.0005631924,0.0004020605,0.00009352678],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.0001479737,0.0001030314,0.009942548,0.0004069812,0.0008483576,0.0001264728,0.008867679,0.0005463835,0.00002193421,0.9659113,0.01217798,0.0008993094],"study_design_scores_gemma":[0.003998466,0.0003336306,0.01136022,0.002337749,0.001297094,0.00003598166,0.01050876,0.07583985,0.00007147052,0.863893,0.0281725,0.002151255],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.7424239,0.006466375,0.0226901,0.002705214,0.003462605,0.001536508,0.0001714258,0.0005538835,0.21999],"genre_scores_gemma":[0.9733674,0.0005062523,0.0002432981,0.0003687954,0.0004223346,0.00001299093,0.0000360025,0.00005724656,0.02498569],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.2309435,"threshold_uncertainty_score":0.9998158,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.07943446247142871,"score_gpt":0.2821092111644805,"score_spread":0.2026747486930518,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}