{"id":"W3171517119","doi":"","title":"LTL2Action: Generalizing LTL Instructions for Multi-Task RL","year":2021,"lang":"en","type":"article","venue":"International Conference on Machine Learning","topic":"Natural Language Processing Techniques","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"ca_institutions":"Vector Institute; University of Toronto","funders":"","keywords":"Computer science; Task (project management); Reinforcement learning; Syntax; Semantics (computer science); Overhead (engineering); Linear temporal logic; Artificial intelligence; Scheme (mathematics); Sample complexity; Programming language; Theoretical computer science; Natural language processing; Mathematics","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001398179,0.0009527731,0.0005501271,0.0003068671,0.000319811,0.0009538821,0.00225125,0.001093403,0.004311731],"category_scores_gemma":[0.004798948,0.0003940719,0.000671536,0.0002829217,0.001358644,0.001718241,0.001896874,0.002599959,0.001104139],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0007491369,"about_ca_system_score_gemma":0.001655515,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.003119265,"about_ca_topic_score_gemma":0.005143838,"domain_scores_codex":[0.9994055,0.0001724775,0.0000460401,0.0001352785,0.0001711891,0.00006958527],"domain_scores_gemma":[0.9986639,0.0005949424,0.0001218531,0.0003425067,0.000174423,0.0001022816],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0002192619,0.0002033914,0.001058137,0.0002089232,0.0000605914,0.0002189685,0.0002909322,0.7392447,0.01139565,0.06233893,0.003573049,0.1811875],"study_design_scores_gemma":[0.00001773496,0.00004237684,0.00004251646,0.00000977856,0.000005615071,0.00001513271,0.000007084656,0.9717156,0.002301726,0.02431438,0.001520864,0.00000723558],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.006582959,0.00005223094,0.9881415,0.0001964159,0.0000468987,0.00008360773,0.00008439182,0.002726268,0.002085673],"genre_scores_gemma":[0.4981858,0.0001709485,0.4943732,0.0006260372,0.00006155716,0.0005845864,0.0004309187,0.0007823563,0.004784563],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.004311731,"threshold_uncertainty_score":0.0144242,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.07083520501108485,"score_gpt":0.3596141055672882,"score_spread":0.2887789005562033,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}