{"id":"W1931877416","doi":"10.1184/r1/6550949","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","year":2010,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Advanced Bandit Algorithms Research","field":"Decision Sciences","cited_by":846,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Office of Naval Research; Multidisciplinary University Research Initiative","keywords":"Regret; Computer science; Benchmark (surveying); Artificial intelligence; Imitation; Online learning; Reduction (mathematics); Machine learning; Sequence (biology); Iterative learning control; Online machine learning; Convergence (economics); Mathematical optimization; Active learning (machine learning); Mathematics; Economics; Psychology","routes":{"ca_aff":false,"ca_fund":true,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002931319,0.001367687,0.00193232,0.000632973,0.0005669019,0.001064072,0.002934896,0.002079901,0.00244464],"category_scores_gemma":[0.01589996,0.0007369491,0.001039364,0.0007185147,0.002182313,0.002346724,0.002802889,0.003679124,0.0006362479],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001367724,"about_ca_system_score_gemma":0.00188769,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.002323702,"about_ca_topic_score_gemma":0.001771048,"domain_scores_codex":[0.9976894,0.0009253963,0.00008111751,0.0005142363,0.0006217405,0.000168072],"domain_scores_gemma":[0.9933974,0.005004345,0.0003969873,0.0006606079,0.0003367552,0.0002038774],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.0001251149,0.0001738742,0.000651872,0.0001451807,0.00007401779,0.0001220534,0.0001323818,0.813426,0.001304635,0.1132639,0.002845851,0.06773507],"study_design_scores_gemma":[0.00001177922,0.00004040041,0.0000675025,0.000007808499,0.000005565972,0.00002388833,0.000004055294,0.9664181,0.0003595644,0.03254486,0.0005103293,0.000006139363],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.00525936,0.0001684598,0.9916065,0.0002916976,0.00004475705,0.00005826696,0.00002541655,0.000269988,0.002275534],"genre_scores_gemma":[0.553184,0.0003842257,0.4371679,0.0005073092,0.0003375725,0.0005328431,0.0002730712,0.0003226741,0.007290335],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.002934896,"threshold_uncertainty_score":0.01550245,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.1356271737258775,"score_gpt":0.3110003567384468,"score_spread":0.1753731830125693,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}