{"id":"W7127642937","doi":"10.1109/cicn67655.2025.11368293","title":"PromptOS: Study of Natural-Language Shells, Execution-Based Evaluation, Command-Syntax Learning, and Runtime Guardrails","year":2025,"lang":"","type":"article","venue":"","topic":"Natural Language Processing Techniques","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"Trinity College","funders":"","keywords":"Undo; Workflow; Cloud computing; Key (lock); Parsing; Overlay; Architecture; Syntax","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008145505,0.001025708,0.0006451673,0.00119768,0.0005864285,0.002938796,0.001986532,0.0008382837,0.003898884],"category_scores_gemma":[0.03964705,0.0006854081,0.000740343,0.0008196793,0.002540305,0.00821838,0.002802755,0.002304114,0.001123195],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001782168,"about_ca_system_score_gemma":0.003125452,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.003515948,"about_ca_topic_score_gemma":0.002558676,"domain_scores_codex":[0.9922092,0.003775249,0.0005213335,0.00107602,0.002104178,0.0003139593],"domain_scores_gemma":[0.9706293,0.01844524,0.001732823,0.005152368,0.003446038,0.0005943049],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0009943983,0.0005529739,0.01768636,0.001146328,0.0001077334,0.0004203556,0.005777234,0.08190114,0.02852018,0.2686663,0.009146936,0.58508],"study_design_scores_gemma":[0.00007944377,0.0007472213,0.004779055,0.0003079291,0.00009267267,0.0006587939,0.001350517,0.7452385,0.1085265,0.09422807,0.04385234,0.0001390654],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.0770787,0.0003020802,0.8967358,0.0004391668,0.00004726663,0.0002399606,0.000269638,0.01834833,0.00653913],"genre_scores_gemma":[0.5039058,0.0003388714,0.4850558,0.000257759,0.00003395529,0.0002463169,0.001218736,0.004140063,0.004802747],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.008145505,"threshold_uncertainty_score":0.04307812,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.00831503984126254,"score_gpt":0.3268386743978949,"score_spread":0.3185236345566324,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}