{"id":"W4409231750","doi":"10.14778/3712221.3712227","title":"How Reliable are Streams? End-to-End Processing-Guarantee Validation and Performance Benchmarking of Stream Processing Systems","year":2024,"lang":"en","type":"article","venue":"Proceedings of the VLDB Endowment","topic":"Distributed systems and fault tolerance","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Toronto","funders":"","keywords":"Benchmarking; Stream processing; STREAMS; Computer science; End-to-end principle; Real-time computing; Distributed computing; Computer network; Business","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01137252,0.000976522,0.0007921665,0.001460127,0.000619518,0.00243584,0.001885622,0.0009825006,0.0009136312],"category_scores_gemma":[0.07249767,0.0004730832,0.000369608,0.001279207,0.001251152,0.003918764,0.001391473,0.001349323,0.0002995686],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001159731,"about_ca_system_score_gemma":0.001577015,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.002991367,"about_ca_topic_score_gemma":0.001879876,"domain_scores_codex":[0.9884597,0.004491101,0.001071168,0.0009960883,0.004322259,0.0006595989],"domain_scores_gemma":[0.9510816,0.02608002,0.00468151,0.008435803,0.008516825,0.001204249],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.002962858,0.0007859865,0.1262127,0.0009874929,0.0004207982,0.0007647815,0.001875282,0.5200634,0.04826692,0.02065074,0.01045706,0.2665519],"study_design_scores_gemma":[0.00004796888,0.0003458694,0.008716609,0.00005987436,0.00004862916,0.0001605705,0.0002618229,0.9369318,0.04128393,0.00948051,0.002615806,0.00004654996],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.6075138,0.0009570543,0.3680466,0.0008865803,0.0001383136,0.0002417893,0.0008947473,0.01796612,0.003354914],"genre_scores_gemma":[0.9570993,0.0001582094,0.04127083,0.00007609788,0.00002240367,0.00008180212,0.0006478784,0.0003751644,0.0002683703],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.01137252,"threshold_uncertainty_score":0.06014436,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.01200811440295121,"score_gpt":0.220075315563754,"score_spread":0.2080672011608028,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}