{"id":"W4407248373","doi":"10.48550/arxiv.2502.04297","title":"Statistical guarantees for continuous-time policy evaluation: blessing of ellipticity and new tradeoffs","year":2025,"lang":"en","type":"preprint","venue":"ArXiv.org","topic":"Agricultural Economics and Policy","field":"Agricultural and Biological Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Blessing; Computer science; Mathematical economics; Economics; Philosophy; Theology","routes":{"ca_aff":false,"ca_fund":true,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02253382,0.0012069,0.001744637,0.00127286,0.0009882654,0.003411054,0.002405638,0.002663792,0.002409242],"category_scores_gemma":[0.1382108,0.001043361,0.0008713009,0.0009321945,0.005506615,0.007084528,0.005214404,0.004853616,0.0003803989],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.002841972,"about_ca_system_score_gemma":0.002900465,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.002613603,"about_ca_topic_score_gemma":0.001201024,"domain_scores_codex":[0.9921084,0.003870027,0.0003913176,0.0008925041,0.002274036,0.0004636085],"domain_scores_gemma":[0.8651206,0.1159541,0.006501504,0.006488911,0.004527517,0.001407262],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.000253131,0.00006349428,0.002050091,0.0001896846,0.00006916555,0.0001449093,0.0002826233,0.7326326,0.002730274,0.2331381,0.0006389625,0.02780697],"study_design_scores_gemma":[0.00000941823,0.00003919872,0.0001671703,0.00003044646,0.000004475526,0.00002410175,0.0000188659,0.9491028,0.0009268444,0.04941714,0.0002471097,0.00001246995],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.02343228,0.0005396618,0.9718125,0.001680257,0.0000493619,0.00002928933,0.00004069517,0.0001593914,0.002256497],"genre_scores_gemma":[0.8374616,0.001086372,0.1574088,0.0005414113,0.0002524885,0.0001665551,0.0001209086,0.0002732998,0.0026886],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.02253382,"threshold_uncertainty_score":0.1191717,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.08367694936171947,"score_gpt":0.3095765126839256,"score_spread":0.2258995633222061,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}