{"id":"W2139515132","doi":"","title":"Learning a value analysis tool for agent evaluation","year":2009,"lang":"en","type":"article","venue":"","topic":"Game Theory and Applications","field":"Decision Sciences","cited_by":18,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Variance (accounting); Estimator; Population; Limit (mathematics); Bellman equation; Function (biology); Domain (mathematical analysis); Monte Carlo method; Value (mathematics); Artificial intelligence; Machine learning; Mathematical optimization; Mathematics; Statistics","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01961254,0.00197087,0.002472337,0.004287614,0.0009651727,0.003443737,0.002819895,0.002621134,0.005712128],"category_scores_gemma":[0.08231912,0.0008838489,0.001366499,0.00191888,0.002644218,0.00637802,0.003484381,0.004384559,0.001456356],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.002891639,"about_ca_system_score_gemma":0.002417381,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00127547,"about_ca_topic_score_gemma":0.001073224,"domain_scores_codex":[0.9890808,0.006225091,0.0006668344,0.00114584,0.002417306,0.000463994],"domain_scores_gemma":[0.9492697,0.04031503,0.002245225,0.002886089,0.00455021,0.0007337433],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0002549168,0.0003079539,0.004576022,0.0002863024,0.0002021252,0.0001680575,0.0003577092,0.3807129,0.001825289,0.250344,0.007133538,0.3538311],"study_design_scores_gemma":[0.00002124269,0.00005595631,0.0001729081,0.00003731734,0.0000145206,0.00003719901,0.00002253065,0.9046646,0.0008944783,0.09290127,0.001162392,0.00001565148],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.001872488,0.00008646229,0.996232,0.0001635637,0.00001658082,0.00006556517,0.00002947407,0.0003826008,0.001151259],"genre_scores_gemma":[0.2563414,0.0002316084,0.7401191,0.0002157579,0.0001260385,0.0007751281,0.0002144225,0.0002346007,0.001741954],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.01961254,"threshold_uncertainty_score":0.1037222,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.178569046960018,"score_gpt":0.4774049186321165,"score_spread":0.2988358716720986,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}