{"id":"W164182584","doi":"","title":"Towards a Trial Plan for Evaluating the COMDAT TD","year":2004,"lang":"en","type":"article","venue":"Defense Technical Information Center (DTIC)","topic":"Human-Automation Interaction and Safety","field":"Psychology","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Computer science; Plan (archaeology); Context (archaeology); Class (philosophy); Operator (biology); Interface (matter); Operations research; Risk analysis (engineering); Artificial intelligence; Engineering; Business; Operating system","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0009725293,0.0001800711,0.00020998,0.0001432648,0.0003045501,0.0001545237,0.0003140289,0.0001781464,0.001574746],"category_scores_gemma":[0.0004876135,0.0001306933,0.0002208631,0.0001669938,0.00008590049,0.0005301962,0.00006571048,0.0003246918,0.001545909],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.000165486,"about_ca_system_score_gemma":0.00006888744,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00003127672,"about_ca_topic_score_gemma":0.00001564798,"domain_scores_codex":[0.9980952,0.0001088846,0.0009408337,0.0001648469,0.0003667746,0.0003234574],"domain_scores_gemma":[0.9986385,0.0002599711,0.0003310392,0.0004629797,0.000213416,0.00009417158],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","study_design_scores_codex":[0.03707553,0.001235783,0.000097857,0.0001092211,0.0003277478,0.000004776323,0.01913362,0.000816432,0.0001889782,0.7069569,0.148984,0.08506917],"study_design_scores_gemma":[0.1100179,0.002704338,0.005243748,0.000134727,0.0001300237,0.0002989336,0.003055676,0.003871582,0.0001921099,0.009205886,0.8643877,0.000757351],"study_design_candidate":"not_applicable","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.2334398,0.0000358832,0.3818392,0.02024893,0.01012966,0.00884401,0.0006502251,0.002300965,0.3425114],"genre_scores_gemma":[0.9901956,0.00000273765,0.002061225,0.006138221,0.0002333384,0.0006448733,0.0004043099,0.00001689637,0.0003028247],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.7567558,"threshold_uncertainty_score":0.999338,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.1168644969045702,"score_gpt":0.4279192292661632,"score_spread":0.311054732361593,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}