{"id":"W6929105423","doi":"10.48448/f1x6-pe65","title":"Human-Machine Teaming Considerations Necessary to Develop Trustworthy, Mission Ready Large Language Models","year":2024,"lang":"en","type":"other","venue":"Open MIND","topic":"","field":"","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Variety (cybernetics); Order (exchange); Work (physics); Trustworthiness; Military intelligence; Silver bullet","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02509171,0.001017535,0.0007523864,0.001696786,0.002939049,0.01116101,0.002635075,0.002178915,0.005136467],"category_scores_gemma":[0.119741,0.001397221,0.001199905,0.00075899,0.004716477,0.01857287,0.008349053,0.005653375,0.002777708],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.002692735,"about_ca_system_score_gemma":0.007137985,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.005812889,"about_ca_topic_score_gemma":0.009732529,"domain_scores_codex":[0.975713,0.015899,0.001321923,0.001894105,0.00456638,0.0006055309],"domain_scores_gemma":[0.9169383,0.05141926,0.003648368,0.01710695,0.008564362,0.002322822],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.0003122701,0.0004244229,0.01205316,0.0009705719,0.0002260119,0.001017846,0.0360866,0.1009119,0.01576426,0.4800977,0.02972251,0.3224126],"study_design_scores_gemma":[0.00006208555,0.0001204354,0.001587553,0.0006048503,0.00007194374,0.0004388923,0.008946977,0.3343906,0.008298791,0.5617662,0.08357386,0.0001378022],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"methods","genre_gemma":"other","genre_scores_codex":[0.02663608,0.0005456793,0.9324837,0.02065669,0.0001770528,0.0005551673,0.0002473599,0.002932851,0.01576558],"genre_scores_gemma":[0.1760546,0.0003177524,0.8189233,0.0007473477,0.0000857406,0.0004247608,0.0004727033,0.0006907279,0.002283034],"genre_candidate":"other","genre_consensus":null,"teacher_disagreement_score":0.02509171,"threshold_uncertainty_score":0.1326992,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.06100700046461459,"score_gpt":0.374165280481201,"score_spread":0.3131582800165864,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}