{"id":"W1493952344","doi":"10.1007/978-3-540-24677-0_102","title":"Multiple Reinforcement Learning Agents in a Static Environment","year":2004,"lang":"en","type":"book","venue":"Lecture notes in computer science","topic":"Reinforcement Learning in Robotics","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"ca_institutions":"Acadia University","funders":"","keywords":"Reinforcement learning; Computer science; Testbed; Reinforcement; Artificial intelligence; Error-driven learning; State (computer science); Learning classifier system; Human–computer interaction; Machine learning; Computer network","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0004629768,0.0007331989,0.0008213432,0.0003808941,0.0007305124,0.0007528246,0.001593338,0.001206963,0.004282507],"category_scores_gemma":[0.001491672,0.0004943327,0.00046358,0.0004074888,0.0007907015,0.001366923,0.001958743,0.0008728693,0.0006756422],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0005731808,"about_ca_system_score_gemma":0.0006647637,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001705184,"about_ca_topic_score_gemma":0.001963001,"domain_scores_codex":[0.9997037,0.00005815486,0.00001454595,0.00007971046,0.0000963104,0.00004754906],"domain_scores_gemma":[0.9994381,0.0002237904,0.00006781613,0.00006331475,0.00008619572,0.0001208837],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0002393696,0.0001536018,0.0006061669,0.0000904284,0.00008139181,0.0004325476,0.0001000661,0.8695744,0.009533966,0.02485191,0.001684578,0.09265153],"study_design_scores_gemma":[0.00003134727,0.00008607509,0.0001724238,0.000006795674,0.00001803982,0.00006533956,0.00002200592,0.984597,0.001582539,0.01187113,0.001536319,0.00001103458],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.1250104,0.000457194,0.8513551,0.0004791122,0.0001809973,0.0001009104,0.00005157065,0.0009358241,0.02142888],"genre_scores_gemma":[0.7827398,0.0004371661,0.1865007,0.0001120601,0.00009520687,0.0001792199,0.00008983402,0.0001057515,0.02974022],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.004282507,"threshold_uncertainty_score":0.01432639,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.02008487822539183,"score_gpt":0.2464408644933316,"score_spread":0.2263559862679398,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}