{"id":"W4402262825","doi":"10.23919/acc60939.2024.10644360","title":"A Practical Reinforcement Learning (RL) Controller Design for Nonlinear Systems","year":2024,"lang":"en","type":"article","venue":"","topic":"Advanced Control Systems Optimization","field":"Engineering","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"McMaster University","funders":"","keywords":"Reinforcement learning; Computer science; Nonlinear system; Controller (irrigation); Reinforcement; Control theory (sociology); Control engineering; Artificial intelligence; Engineering; Control (management); Structural engineering","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009029756,0.0008767787,0.0006226615,0.0002108425,0.0003089863,0.0007027006,0.000832524,0.0009331501,0.002307191],"category_scores_gemma":[0.001514197,0.0003204953,0.0004248352,0.0001652459,0.0007318786,0.0004945379,0.0008346444,0.001191855,0.0005516949],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0004309098,"about_ca_system_score_gemma":0.0009877792,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.002670935,"about_ca_topic_score_gemma":0.001924588,"domain_scores_codex":[0.9995466,0.0001200051,0.00002274317,0.0001024627,0.0001715734,0.00003660317],"domain_scores_gemma":[0.9995853,0.000173172,0.00005916554,0.00003192024,0.0001315143,0.00001879362],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.00005548215,0.0000469927,0.0002007124,0.0002056907,0.00002693576,0.0001188739,0.00007942814,0.9184298,0.009309163,0.01064245,0.00107072,0.05981373],"study_design_scores_gemma":[0.00001420956,0.00006135321,0.0000384512,0.000008383619,0.000004511578,0.00001978757,0.000003520167,0.9970015,0.0008380262,0.00090003,0.001106058,0.000004163202],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.002302343,0.0001592374,0.9945643,0.0001209731,0.00003974799,0.00005155648,0.000008393247,0.0002052801,0.002548098],"genre_scores_gemma":[0.7838038,0.0004794209,0.2088703,0.000218693,0.0001072841,0.0004936472,0.00006373329,0.00007081839,0.005892343],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.002670935,"threshold_uncertainty_score":0.007718265,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.02457321598695149,"score_gpt":0.2777408226012688,"score_spread":0.2531676066143173,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}