{"id":"W4385718290","doi":"10.21203/rs.3.rs-3221401/v1","title":"Smart Data Prefetching Using KNN to Improve Hadoop Performance","year":2023,"lang":"en","type":"preprint","venue":"Research Square","topic":"Cloud Computing and Resource Management","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"McMaster University","funders":"","keywords":"Instruction prefetch; Computer science; Cluster analysis; Locality; Real-time computing; Distributed computing; Operating system; Cache; Machine learning","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001175028,0.0008792832,0.001004345,0.001260262,0.001413859,0.000748957,0.001778923,0.0004862141,0.001183917],"category_scores_gemma":[0.003979865,0.0004726895,0.0003906288,0.001698373,0.0006602681,0.001865509,0.001011455,0.0007345673,0.0004395996],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001114848,"about_ca_system_score_gemma":0.002166028,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.01305789,"about_ca_topic_score_gemma":0.02046661,"domain_scores_codex":[0.9990194,0.0001405354,0.00008517571,0.0002366116,0.0003045005,0.00021379],"domain_scores_gemma":[0.9973916,0.0007716613,0.0002609914,0.0004379771,0.0008521704,0.0002856025],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.002070175,0.000867741,0.01965756,0.0007398095,0.0002163896,0.0005038504,0.0004815922,0.5453932,0.09389625,0.003698393,0.01600428,0.3164707],"study_design_scores_gemma":[0.00005542564,0.0001387059,0.001166644,0.00001695899,0.000027883,0.00006195012,0.00009509345,0.9792044,0.01469803,0.002437392,0.002059346,0.0000381106],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.4537992,0.003524139,0.513395,0.000931197,0.00087144,0.0004396271,0.0006595859,0.02018281,0.006197078],"genre_scores_gemma":[0.8647773,0.0002519285,0.1331829,0.0001352576,0.00007376556,0.00007690681,0.0003215003,0.0001976603,0.0009827659],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.01305789,"threshold_uncertainty_score":0.02596378,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.2371251602670578,"score_gpt":0.4206197613968451,"score_spread":0.1834946011297873,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}