{"id":"W1886179201","doi":"10.1007/11424918_5","title":"ARES 2: A Tool for Evaluating Cooperative and Competitive Multi-agent Systems","year":2005,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Multi-Agent Systems and Negotiation","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Key (lock); Systems engineering; Range (aeronautics); Multi-agent system; Software engineering; Simulation; Operations research; Artificial intelligence; Computer security; Aerospace engineering; Engineering","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaepi_narrow"],"consensus_categories":[],"category_scores_codex":[0.001559545,0.0005416531,0.0006538802,0.0004896228,0.0004441504,0.0009018194,0.001225528,0.0002628204,0.000007408708],"category_scores_gemma":[0.0002102673,0.0004717401,0.0001069995,0.0002496501,0.0003236071,0.0006499542,0.0006399527,0.0003857908,0.00001768673],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0003728315,"about_ca_system_score_gemma":0.0003816726,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00003740886,"about_ca_topic_score_gemma":0.00009993293,"domain_scores_codex":[0.9962195,0.0000930029,0.0007116591,0.001579866,0.0008325118,0.00056351],"domain_scores_gemma":[0.9971526,0.0008937497,0.0004653754,0.0007453397,0.0006012133,0.0001417033],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.00003121888,0.000127722,0.0002679637,0.0004758209,0.000096245,0.00004304406,0.00802989,0.220079,0.00148233,0.1811291,0.00006339819,0.5881743],"study_design_scores_gemma":[0.0007881793,0.0002826087,0.0002641874,0.0008465855,0.00001320738,0.00004381418,0.000001695869,0.9938772,0.0003876401,0.001325726,0.001557043,0.0006121228],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.0003700365,0.001315348,0.9938012,0.0002526509,0.001870453,0.001966398,0.00002969602,0.0001103593,0.0002838054],"genre_scores_gemma":[0.2947887,0.00008848718,0.701754,0.0008126494,0.001029663,0.0001730823,0.00002046661,0.00005979019,0.001273224],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.7737982,"threshold_uncertainty_score":0.9997734,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.04848181622520149,"score_gpt":0.3061279779522542,"score_spread":0.2576461617270527,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}