{"id":"W2138857742","doi":"","title":"Why Does Unsupervised Pre-training Help Deep Learning?","year":2010,"lang":"en","type":"article","venue":"","topic":"Domain Adaptation and Few-Shot Learning","field":"Computer Science","cited_by":2115,"is_retracted":false,"has_abstract":true,"ca_institutions":"Université de Montréal","funders":"","keywords":"Artificial intelligence; Unsupervised learning; Computer science; Machine learning; Deep learning; Regularization (linguistics); Generalization; Autoencoder; Deep belief network; Semi-supervised learning; Competitive learning; Training (meteorology); Mathematics","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007312512,0.001713834,0.001877332,0.0005122933,0.0007798735,0.001805452,0.00230202,0.003385202,0.003870591],"category_scores_gemma":[0.03210307,0.0008038565,0.0007319591,0.0007758,0.002538531,0.006976046,0.002159683,0.005198376,0.003072922],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0006067057,"about_ca_system_score_gemma":0.001190027,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001503274,"about_ca_topic_score_gemma":0.002761448,"domain_scores_codex":[0.9976816,0.001162611,0.000103538,0.0005517757,0.0003195066,0.000180923],"domain_scores_gemma":[0.9823805,0.01153116,0.0006536071,0.003406957,0.001511196,0.0005166743],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.001026721,0.0007304428,0.01117721,0.001463766,0.0004199435,0.0002769861,0.0006978896,0.07441362,0.02356208,0.06667838,0.03700478,0.7825482],"study_design_scores_gemma":[0.0001767607,0.0006973048,0.006342016,0.0007226206,0.0001684217,0.0006857694,0.0005993724,0.5636077,0.03956596,0.3532561,0.0339977,0.0001802951],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.04026416,0.006137257,0.9225798,0.01810826,0.0007433931,0.0001453602,0.0002358403,0.003706368,0.008079656],"genre_scores_gemma":[0.4773853,0.005274984,0.4993584,0.005797181,0.001245617,0.0002673788,0.0006862874,0.001355825,0.008629015],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.007312512,"threshold_uncertainty_score":0.03867269,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.01361732876810003,"score_gpt":0.2434216545242386,"score_spread":0.2298043257561385,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}