{"id":"W2950296358","doi":"10.1088/1742-6596/1213/2/022021","title":"Research on data cleaning technology based on instance level","year":2019,"lang":"en","type":"article","venue":"Journal of Physics Conference Series","topic":"Data Quality and Management","field":"Decision Sciences","cited_by":4,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Guelph","funders":"","keywords":"Computer science; Process (computing); Data quality; Quality (philosophy); Data mining; Engineering; Programming language","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00373189,0.0005499762,0.001252389,0.003187933,0.001039666,0.004403795,0.002728168,0.001136766,0.001989291],"category_scores_gemma":[0.01082275,0.0005537186,0.001999881,0.004029166,0.001371807,0.007673438,0.001902073,0.002964004,0.0006220746],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001171024,"about_ca_system_score_gemma":0.00167744,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001515138,"about_ca_topic_score_gemma":0.0005590831,"domain_scores_codex":[0.9931398,0.001468302,0.0005572797,0.001673667,0.002838295,0.0003226456],"domain_scores_gemma":[0.9924257,0.002534218,0.0007120294,0.002157688,0.001993598,0.000176763],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0004423015,0.0002850784,0.01483799,0.001458456,0.0005467811,0.0003495458,0.0007722407,0.03433707,0.07496901,0.1755219,0.008887347,0.6875923],"study_design_scores_gemma":[0.00007239066,0.0004323806,0.00814866,0.000309204,0.0004969956,0.001386903,0.0006575459,0.5648112,0.2175642,0.1475404,0.05839004,0.0001900376],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.01659392,0.002215664,0.9766328,0.0008875744,0.000168109,0.00008944871,0.0001874898,0.001224421,0.002000597],"genre_scores_gemma":[0.3642964,0.004303972,0.6265012,0.0007019819,0.0003157896,0.0001879968,0.001059322,0.0002989864,0.002334324],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.004403795,"threshold_uncertainty_score":0.01973635,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.6678248949701602,"score_gpt":0.519995239498033,"score_spread":0.1478296554721272,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}