{"id":"W2569695769","doi":"10.1007/978-3-319-50502-2_3","title":"Markov Multi-armed Bandit","year":2016,"lang":"en","type":"book-chapter","venue":"Wireless networks","topic":"Advanced Bandit Algorithms Research","field":"Decision Sciences","cited_by":2,"is_retracted":false,"has_abstract":false,"ca_institutions":"McMaster University","funders":"","keywords":"Markov chain; Computer science; Markov model; Variable-order Markov model; Markov process; Markov property; Markov decision process; Markov kernel; Mathematical optimization; Markov renewal process; Examples of Markov chains; Mathematics; Machine learning; Statistics","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000827648,0.00127404,0.001016531,0.0005054727,0.0004108664,0.0029311,0.001009263,0.00157743,0.01163968],"category_scores_gemma":[0.003078763,0.0005295951,0.0005114677,0.001505276,0.001065638,0.00198903,0.0009196544,0.002317021,0.004853975],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0009889416,"about_ca_system_score_gemma":0.0007850174,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001123218,"about_ca_topic_score_gemma":0.001482549,"domain_scores_codex":[0.999544,0.0001532832,0.00002231713,0.0000761359,0.000163886,0.00004040903],"domain_scores_gemma":[0.9989358,0.000800323,0.00004672736,0.00009833772,0.00009201221,0.0000268475],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.00004723358,0.00005831801,0.0001577325,0.0002549547,0.00005509546,0.00006105185,0.00008445034,0.0970147,0.0008089372,0.7016353,0.04434926,0.155473],"study_design_scores_gemma":[0.00001544381,0.00002563315,0.0001364741,0.0001390199,0.00002092871,0.00008331358,0.00003060178,0.26733,0.0004750962,0.6729172,0.05879718,0.00002909814],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.00241014,0.01374242,0.8339235,0.001975307,0.001077069,0.00003939622,0.0002224614,0.0004129803,0.1461968],"genre_scores_gemma":[0.2783452,0.05435797,0.3224498,0.002348518,0.004037116,0.000576729,0.001198954,0.0007333423,0.3359523],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.01163968,"threshold_uncertainty_score":0.03893864,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.1022131276347232,"score_gpt":0.3813151395114032,"score_spread":0.27910201187668,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}