{"id":"W1603884869","doi":"10.1007/978-3-540-77990-2_7","title":"On the Evaluation of Agent-Oriented Software Engineering Methodologies: A Statistical Approach","year":2008,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Multi-Agent Systems and Negotiation","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"ca_institutions":"University of Calgary; University of Alberta","funders":"","keywords":"Agent-oriented software engineering; Computer science; Software engineering; Software; Set (abstract data type); Software development; Systems engineering; Artificial intelligence; Engineering","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003385781,0.0003423349,0.0004154817,0.0004574593,0.0001638242,0.00009902176,0.00156231,0.0002045564,0.00001724445],"category_scores_gemma":[0.001943432,0.0002399772,0.00009604907,0.0004463223,0.0002718781,0.0002038197,0.0004430189,0.000476387,0.00001037859],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0002898734,"about_ca_system_score_gemma":0.0003578404,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00001806382,"about_ca_topic_score_gemma":0.000003496383,"domain_scores_codex":[0.9961008,0.0002036136,0.0005195772,0.0009334741,0.001887982,0.0003545749],"domain_scores_gemma":[0.9955127,0.002595783,0.0003437513,0.001075003,0.0004004939,0.00007233385],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0000079621,0.00005142548,0.00003446233,0.00009215337,0.00003315566,0.00001210338,0.00245617,0.6077673,0.0001228227,0.08129479,0.0001306288,0.307997],"study_design_scores_gemma":[0.0002057376,0.00009711637,0.0005673725,0.0001970756,0.0000128429,0.0000211283,3.479679e-7,0.9886066,0.0004808979,0.009402725,0.0001467206,0.000261476],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.0003488064,0.000208158,0.9971054,0.00008035024,0.001116733,0.0007117928,0.000008890253,0.0000891018,0.0003307864],"genre_scores_gemma":[0.1383962,0.000021325,0.8611478,0.0001998229,0.0001392217,0.00003146294,0.00001055719,0.00002058368,0.00003303043],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.3808393,"threshold_uncertainty_score":0.9785984,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.1146426977546878,"score_gpt":0.3170349928897275,"score_spread":0.2023922951350398,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}