{"id":"W3178495416","doi":"10.48550/arxiv.2107.06196","title":"No Regrets for Learning the Prior in Bandits","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Advanced Bandit Algorithms Research","field":"Decision Sciences","cited_by":1,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Alberta","funders":"","keywords":"Regret; Task (project management); Thompson sampling; Key (lock); Computer science; Bayesian probability; Bayes' theorem; Artificial intelligence; Sampling (signal processing); Prior probability; Machine learning; Mathematical optimization; Mathematics; Engineering; Computer vision","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007806905,0.002039906,0.002623146,0.000767905,0.001650948,0.002660201,0.002867002,0.003235589,0.004549328],"category_scores_gemma":[0.04247747,0.001060354,0.001150446,0.001028946,0.003218669,0.00462378,0.002915897,0.004758613,0.00119825],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.00268979,"about_ca_system_score_gemma":0.002591548,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.003907772,"about_ca_topic_score_gemma":0.00404845,"domain_scores_codex":[0.9943795,0.003344019,0.0001946378,0.000816918,0.0008268034,0.0004381587],"domain_scores_gemma":[0.9761594,0.01931112,0.001035954,0.001983759,0.0008392301,0.0006705705],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.0006255785,0.0002159504,0.001683924,0.000198498,0.0001126288,0.0000999836,0.000203377,0.7683005,0.001259538,0.1597065,0.006187513,0.06140596],"study_design_scores_gemma":[0.0000408066,0.00004705504,0.0001398301,0.0000216482,0.00001296571,0.00001978805,0.00001381412,0.9059085,0.0003397919,0.09283816,0.000605922,0.00001164321],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.03835438,0.001024668,0.9504519,0.001653906,0.0001496126,0.0001102299,0.0001570884,0.0007906447,0.007307559],"genre_scores_gemma":[0.7638656,0.0007674488,0.22297,0.001161066,0.0003460674,0.0006599753,0.0004907385,0.0005809026,0.009158147],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.007806905,"threshold_uncertainty_score":0.04128736,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.2585587526094204,"score_gpt":0.3217372327608468,"score_spread":0.06317848015142646,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}