{"id":"W4389519544","doi":"10.18653/v1/2023.newsum-1.12","title":"Analyzing Multi-Sentence Aggregation in Abstractive Summarization via the Shapley Value","year":2023,"lang":"en","type":"article","venue":"","topic":"Advanced Text Analysis Techniques","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"ca_institutions":"","funders":"Canadian Institute for Advanced Research; Nvidia","keywords":"Automatic summarization; Shapley value; Computer science; Sentence; Value (mathematics); Natural language processing; Artificial intelligence; Mathematics; Game theory; Machine learning; Mathematical economics","routes":{"ca_aff":false,"ca_fund":true,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006163673,0.0008606435,0.001954226,0.003020543,0.001371825,0.003581522,0.001474792,0.001274853,0.003831856],"category_scores_gemma":[0.02689144,0.0005124105,0.001095679,0.003280399,0.001362138,0.006756174,0.001918241,0.00185031,0.0004735364],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001490735,"about_ca_system_score_gemma":0.001415799,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001440954,"about_ca_topic_score_gemma":0.00173575,"domain_scores_codex":[0.9958028,0.002188905,0.0002444597,0.0005994581,0.0009003361,0.0002641551],"domain_scores_gemma":[0.977815,0.01751672,0.001041716,0.001211747,0.001972523,0.0004423328],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"bench_or_experimental","study_design_scores_codex":[0.0006300327,0.0004394092,0.004203546,0.000750754,0.000687686,0.0003119238,0.001183584,0.2925249,0.006455566,0.4189008,0.009547352,0.2643645],"study_design_scores_gemma":[0.00002711736,0.00008766113,0.0007840359,0.00003518466,0.00007290374,0.00003413305,0.0001076027,0.6119676,0.001129063,0.3846738,0.001056895,0.0000238755],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.1153678,0.001068455,0.8745943,0.0009369505,0.0001374983,0.0001888307,0.000396987,0.0003714879,0.006937655],"genre_scores_gemma":[0.8351517,0.0004883707,0.1584019,0.00022516,0.0003186033,0.000236755,0.0009671638,0.0001794677,0.004030918],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.006163673,"threshold_uncertainty_score":0.03259701,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.02542860140765619,"score_gpt":0.3057754320863267,"score_spread":0.2803468306786705,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}