{"id":"W4320342078","doi":"10.48550/arxiv.2302.00237","title":"Bridging Physics-Informed Neural Networks with Reinforcement Learning: Hamilton-Jacobi-Bellman Proximal Policy Optimization (HJBPPO)","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Reinforcement Learning in Robotics","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"University of Waterloo","keywords":"Hamilton–Jacobi–Bellman equation; Reinforcement learning; Bellman equation; Artificial neural network; Bridging (networking); Computer science; Mathematical optimization; Optimal control; Mathematics; Artificial intelligence","routes":{"ca_aff":false,"ca_fund":true,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001894337,0.0008374127,0.001090939,0.0003718048,0.0005166814,0.0009616132,0.001411233,0.001566327,0.00207735],"category_scores_gemma":[0.005385756,0.0006204315,0.0004414595,0.0004944137,0.001723813,0.0013331,0.002070179,0.002284233,0.0003915562],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001006578,"about_ca_system_score_gemma":0.00209271,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.004460794,"about_ca_topic_score_gemma":0.003006794,"domain_scores_codex":[0.9993601,0.0003213409,0.00002128068,0.00008847204,0.0001558896,0.00005280408],"domain_scores_gemma":[0.9987545,0.0008384524,0.00009779898,0.0001027299,0.0001379408,0.0000685144],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.00004410854,0.00004332703,0.0002851064,0.00005793364,0.00003221278,0.00005109736,0.00005354094,0.9046214,0.0007618485,0.0632538,0.0009749522,0.02982067],"study_design_scores_gemma":[0.000007396123,0.00001487918,0.00002651418,0.000005422531,0.000002586679,0.000007277812,0.000002824912,0.9815224,0.0002421658,0.01769031,0.0004738338,0.000004386339],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.004594462,0.0001752135,0.9927614,0.0002266179,0.00004246098,0.00002418952,0.000008494561,0.0001385602,0.002028495],"genre_scores_gemma":[0.5702636,0.0005241221,0.4232805,0.0004738021,0.0001199569,0.0002555607,0.00006375366,0.0001719731,0.004846611],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.004460794,"threshold_uncertainty_score":0.01001835,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.05811504637595478,"score_gpt":0.2070428115423661,"score_spread":0.1489277651664113,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}