{"id":"W4408534506","doi":"10.3390/a18030170","title":"Investigations into the Design and Implementation of Reinforcement Learning Using Deep Learning Neural Networks","year":2025,"lang":"en","type":"article","venue":"Algorithms","topic":"Advanced Research in Systems and Signal Processing","field":"Engineering","cited_by":1,"is_retracted":false,"has_abstract":true,"ca_institutions":"John Abbott College; Concordia University","funders":"","keywords":"Reinforcement learning; Computer science; Artificial neural network; Artificial intelligence; Deep learning; Machine learning; Deep neural networks","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00129596,0.0004718974,0.0003402546,0.0001738787,0.0002025091,0.0005847982,0.0009023452,0.0006711467,0.001043417],"category_scores_gemma":[0.002799155,0.0002631555,0.0002445886,0.0001159327,0.0006483244,0.0005530954,0.0003367407,0.001011499,0.0001562098],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0007389301,"about_ca_system_score_gemma":0.0009454573,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.003395926,"about_ca_topic_score_gemma":0.00299448,"domain_scores_codex":[0.9996488,0.00009864893,0.00002172165,0.00005811923,0.0001329554,0.00003973212],"domain_scores_gemma":[0.9989775,0.0005297573,0.0001183309,0.00008871939,0.0002489008,0.00003672587],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.00005247275,0.00007910517,0.0007227975,0.0001130835,0.00002741119,0.00004991081,0.00006280345,0.941374,0.005810237,0.009774826,0.0002223162,0.04171103],"study_design_scores_gemma":[0.000007569873,0.00005681927,0.00007374858,0.000007386953,0.000003630752,0.00001024976,0.00000396562,0.9968045,0.001887233,0.0006191185,0.0005233334,0.000002505175],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.041145,0.0003160033,0.9521609,0.0002680208,0.00005885794,0.0001319345,0.00001396264,0.0003183363,0.005586995],"genre_scores_gemma":[0.8909075,0.0002032779,0.1070307,0.00009978622,0.00001846352,0.0001067813,0.00001455512,0.00002626419,0.001592593],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.003395926,"threshold_uncertainty_score":0.006853759,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.02188067409461759,"score_gpt":0.3093974119279932,"score_spread":0.2875167378333756,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}