{"id":"W1603884869","doi":"10.1007/978-3-540-77990-2_7","title":"On the Evaluation of Agent-Oriented Software Engineering Methodologies: A Statistical Approach","year":2008,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Multi-Agent Systems and Negotiation","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"ca_institutions":"University of Calgary; University of Alberta","funders":"","keywords":"Agent-oriented software engineering; Computer science; Software engineering; Software; Set (abstract data type); Software development; Systems engineering; Artificial intelligence; Engineering","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.06484749,0.00140563,0.002515284,0.006968886,0.001174354,0.005871981,0.002961522,0.002065846,0.002548276],"category_scores_gemma":[0.2665772,0.0007356054,0.001477963,0.005258702,0.004578752,0.007318638,0.003265794,0.002395686,0.0003362345],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.003122452,"about_ca_system_score_gemma":0.003709097,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.002403478,"about_ca_topic_score_gemma":0.001899698,"domain_scores_codex":[0.9294867,0.04985958,0.002229398,0.002394173,0.01513995,0.0008902462],"domain_scores_gemma":[0.5472104,0.4170924,0.009261895,0.01270939,0.01244538,0.00128042],"domain_codex":null,"domain_gemma":"evaluation","domain_candidate":"evaluation","domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","study_design_scores_codex":[0.000403978,0.0003749729,0.01539595,0.0006107657,0.0006235822,0.0001385257,0.0007327022,0.2477858,0.001412071,0.4988366,0.00361146,0.2300736],"study_design_scores_gemma":[0.00004579076,0.0003187619,0.003941578,0.000187717,0.0001042659,0.00007625497,0.0002519986,0.641906,0.001282546,0.3494853,0.002355873,0.00004387306],"study_design_candidate":"observational","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.04155054,0.002848851,0.94254,0.001347167,0.0001257742,0.0002780784,0.0001679419,0.0003694439,0.01077208],"genre_scores_gemma":[0.6916407,0.002226189,0.3016803,0.0003168847,0.0004218361,0.0009834147,0.0005051889,0.000304936,0.001920629],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.9351525,"threshold_uncertainty_score":0.3429504,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.1146426977546878,"score_gpt":0.3170349928897275,"score_spread":0.2023922951350398,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}