release 2014.11.26.2

release 2014.11.26.1
[test_unicode_literals] Arm unicode_literals check
2014-11-26 22:06:27 +01:00 · 2014-11-26 22:01:06 +01:00 · 2014-11-26 20:01:22 +01:00 · 2014-11-27 00:24:19 +06:00 · 2014-11-26 21:25:43 +06:00 · 2014-11-26 21:02:46 +06:00
286 changed files with 3862 additions and 2320 deletions
--- a/8
+++ b/8
@ -80,3 +80,11 @@ Damon Timm
 winwon
 Xavier Beynon
 Gabriel Schubiner
 xantares
 Jan Matějka
 Mauroy Sébastien
 William Sewell
 Dao Hoang Son
 Oskar Jauch
 Matthew Rayfield
 t0mm0
--- a/README.md
+++ b/README.md
@ -30,7 +30,7 @@ Alternatively, refer to the developer instructions below for how to check out an
 # DESCRIPTION
 **youtube-dl** is a small command-line program to download videos from
 YouTube.com and a few more sites. It requires the Python interpreter, version
-2.6, 2.7, or 3.3+, and it is not platform specific. It should work on
+2.6, 2.7, or 3.2+, and it is not platform specific. It should work on
 your Unix box, on Windows or on Mac OS X. It is released to the public domain,
 which means you can modify it, redistribute it or use it however you like.
@ -93,7 +93,8 @@ which means you can modify it, redistribute it or use it however you like.
                                     COUNT views
    --max-views COUNT                Do not download any videos with more than
                                     COUNT views
-    --no-playlist                    download only the currently playing video
+    --no-playlist                    If the URL refers to a video and a
                                     playlist, download only the video.
    --age-limit YEARS                download only videos suitable for the given
                                     age
    --download-archive FILE          Download only videos not listed in the
@ -131,17 +132,19 @@ which means you can modify it, redistribute it or use it however you like.
                                     %(upload_date)s for the upload date
                                     (YYYYMMDD), %(extractor)s for the provider
                                     (youtube, metacafe, etc), %(id)s for the
-                                     video id, %(playlist)s for the playlist the
+                                     video id, %(playlist_title)s,
                                     %(playlist_id)s, or %(playlist)s (=title if
                                     present, ID otherwise) for the playlist the
                                     video is in, %(playlist_index)s for the
-                                     position in the playlist and %% for a
+                                     position in the playlist. %(height)s and
-                                     literal percent. %(height)s and %(width)s
+                                     %(width)s for the width and height of the
-                                     for the width and height of the video
+                                     video format. %(resolution)s for a textual
                                     format. %(resolution)s for a textual
                                     description of the resolution of the video
-                                     format. Use - to output to stdout. Can also
+                                     format. %% for a literal percent. Use - to
-                                     be used to download to a different
+                                     output to stdout. Can also be used to
-                                     directory, for example with -o '/my/downloa
+                                     download to a different directory, for
-                                     ds/%(uploader)s/%(title)s-%(id)s.%(ext)s' .
+                                     example with -o '/my/downloads/%(uploader)s
                                     /%(title)s-%(id)s.%(ext)s' .
    --autonumber-size NUMBER         Specifies the number of digits in
                                     %(autonumber)s when it is present in output
                                     filename template or --auto-number option
@ -240,7 +243,12 @@ which means you can modify it, redistribute it or use it however you like.
                                     default, youtube-dl will pick the best
                                     quality. Use commas to download multiple
                                     audio formats, such as -f
-                                     136/137/mp4/bestvideo,140/m4a/bestaudio
+                                     136/137/mp4/bestvideo,140/m4a/bestaudio.
                                     You can merge the video and audio of two
                                     formats into a single file using -f <video-
                                     format>+<audio-format> (requires ffmpeg or
                                     avconv), for example -f
                                     bestvideo+bestaudio.
    --all-formats                    download all available video formats
    --prefer-free-formats            prefer free video formats unless a specific
                                     one is requested
@ -485,14 +493,15 @@ If you want to add support for a new site, you can follow this quick list (assum
        def _real_extract(self, url):
            video_id = self._match_id(url)
            webpage = self._download_webpage(url, video_id)
            # TODO more code goes here, for example ...
            webpage = self._download_webpage(url, video_id)
            title = self._html_search_regex(r'<h1>(.*?)</h1>', webpage, 'title')
            return {
                'id': video_id,
                'title': title,
                'description': self._og_search_description(webpage),
                # TODO more properties (see youtube_dl/extractor/common.py)
            }
    ```
@ -500,7 +509,7 @@ If you want to add support for a new site, you can follow this quick list (assum
 6. Run `python test/test_download.py TestDownload.test_YourExtractor`. This *should fail* at first, but you can continually re-run it until you're done. If you decide to add more than one test, then rename ``_TEST`` to ``_TESTS`` and make it into a list of dictionaries. The tests will be then be named `TestDownload.test_YourExtractor`, `TestDownload.test_YourExtractor_1`, `TestDownload.test_YourExtractor_2`, etc.
 7. Have a look at [`youtube_dl/common/extractor/common.py`](https://github.com/rg3/youtube-dl/blob/master/youtube_dl/extractor/common.py) for possible helper methods and a [detailed description of what your extractor should return](https://github.com/rg3/youtube-dl/blob/master/youtube_dl/extractor/common.py#L38). Add tests and code for as many as you want.
 8. If you can, check the code with [pyflakes](https://pypi.python.org/pypi/pyflakes) (a good idea) and [pep8](https://pypi.python.org/pypi/pep8) (optional, ignore E501).
-9. When the tests pass, [add](https://www.kernel.org/pub/software/scm/git/docs/git-add.html) the new files and [commit](https://www.kernel.org/pub/software/scm/git/docs/git-commit.html) them and [push](https://www.kernel.org/pub/software/scm/git/docs/git-push.html) the result, like this:
+9. When the tests pass, [add](http://git-scm.com/docs/git-add) the new files and [commit](http://git-scm.com/docs/git-commit) them and [push](http://git-scm.com/docs/git-push) the result, like this:
        $ git add youtube_dl/extractor/__init__.py
        $ git add youtube_dl/extractor/yourextractor.py
--- a/devscripts/bash-completion.py
+++ b/devscripts/bash-completion.py
@ -1,4 +1,6 @@
 #!/usr/bin/env python
 from __future__ import unicode_literals
 import os
 from os.path import dirname as dirn
 import sys
@ -9,6 +11,7 @@ import youtube_dl
 BASH_COMPLETION_FILE = "youtube-dl.bash-completion"
 BASH_COMPLETION_TEMPLATE = "devscripts/bash-completion.in"
 def build_completion(opt_parser):
    opts_flag = []
    for group in opt_parser.option_groups:
--- a/devscripts/buildserver.py
+++ b/devscripts/buildserver.py
@ -233,6 +233,7 @@ def rmtree(path):
 #==============================================================================
 class BuildError(Exception):
    def __init__(self, output, code=500):
        self.output = output
--- a/devscripts/check-porn.py
+++ b/devscripts/check-porn.py
@ -1,4 +1,5 @@
 #!/usr/bin/env python
 from __future__ import unicode_literals
 """
 This script employs a VERY basic heuristic ('porn' in webpage.lower()) to check
--- a/devscripts/fish-completion.py
+++ b/devscripts/fish-completion.py
@ -23,13 +23,13 @@ EXTRA_ARGS = {
    'batch-file': ['--require-parameter'],
 }
 def build_completion(opt_parser):
    commands = []
    for group in opt_parser.option_groups:
        for option in group.option_list:
            long_option = option.get_opt_string().strip('-')
            help_msg = shell_quote([option.help])
            complete_cmd = ['complete', '--command', 'youtube-dl', '--long-option', long_option]
            if option._short_opts:
                complete_cmd += ['--short-option', option._short_opts[0].strip('-')]
--- a/devscripts/gh-pages/add-version.py
+++ b/devscripts/gh-pages/add-version.py
@ -1,4 +1,5 @@
 #!/usr/bin/env python3
 from __future__ import unicode_literals
 import json
 import sys
--- a/devscripts/gh-pages/generate-download.py
+++ b/devscripts/gh-pages/generate-download.py
@ -1,8 +1,7 @@
 #!/usr/bin/env python3
 from __future__ import unicode_literals
 import hashlib
 import shutil
 import subprocess
 import tempfile
 import urllib.request
 import json
--- a/devscripts/gh-pages/sign-versions.py
+++ b/devscripts/gh-pages/sign-versions.py
@ -1,4 +1,5 @@
 #!/usr/bin/env python3
 from __future__ import unicode_literals, with_statement
 import rsa
 import json
@ -29,4 +30,5 @@ signature = hexlify(rsa.pkcs1.sign(json.dumps(versions_info, sort_keys=True).enc
 print('signature: ' + signature)
 versions_info['signature'] = signature
-json.dump(versions_info, open('update/versions.json', 'w'), indent=4, sort_keys=True)
+with open('update/versions.json', 'w') as versionsf:
    json.dump(versions_info, versionsf, indent=4, sort_keys=True)
--- a/devscripts/gh-pages/update-copyright.py
+++ b/devscripts/gh-pages/update-copyright.py
@ -1,7 +1,7 @@
 #!/usr/bin/env python
 # coding: utf-8
-from __future__ import with_statement
+from __future__ import with_statement, unicode_literals
 import datetime
 import glob
@ -13,7 +13,7 @@ year = str(datetime.datetime.now().year)
 for fn in glob.glob('*.html*'):
    with io.open(fn, encoding='utf-8') as f:
        content = f.read()
-    newc = re.sub(u'(?P<copyright>Copyright © 2006-)(?P<year>[0-9]{4})', u'Copyright © 2006-' + year, content)
+    newc = re.sub(r'(?P<copyright>Copyright © 2006-)(?P<year>[0-9]{4})', 'Copyright © 2006-' + year, content)
    if content != newc:
        tmpFn = fn + '.part'
        with io.open(tmpFn, 'wt', encoding='utf-8') as outf:
--- a/devscripts/gh-pages/update-feed.py
+++ b/devscripts/gh-pages/update-feed.py
@ -1,4 +1,5 @@
 #!/usr/bin/env python3
 from __future__ import unicode_literals
 import datetime
 import io
@ -73,4 +74,3 @@ atom_template = atom_template.replace('@ENTRIES@', entries_str)
 with io.open('update/releases.atom', 'w', encoding='utf-8') as atom_file:
    atom_file.write(atom_template)
--- a/devscripts/gh-pages/update-sites.py
+++ b/devscripts/gh-pages/update-sites.py
@ -1,4 +1,5 @@
 #!/usr/bin/env python3
 from __future__ import unicode_literals
 import sys
 import os
@ -9,6 +10,7 @@ sys.path.append(os.path.dirname(os.path.dirname(os.path.dirname(os.path.abspath(
 import youtube_dl
 def main():
    with open('supportedsites.html.in', 'r', encoding='utf-8') as tmplf:
        template = tmplf.read()
@ -21,7 +23,7 @@ def main():
            continue
        elif ie_desc is not None:
            ie_html += ': {}'.format(ie.IE_DESC)
-        if ie.working() == False:
+        if not ie.working():
            ie_html += ' (Currently broken)'
        ie_htmls.append('<li>{}</li>'.format(ie_html))
--- a/devscripts/make_readme.py
+++ b/devscripts/make_readme.py
@ -1,3 +1,5 @@
 from __future__ import unicode_literals
 import io
 import sys
 import re
--- a/devscripts/prepare_manpage.py
+++ b/devscripts/prepare_manpage.py
@ -1,3 +1,4 @@
 from __future__ import unicode_literals
 import io
 import os.path
--- a/devscripts/transition_helper.py
+++ b/devscripts/transition_helper.py
@ -1,40 +0,0 @@
 #!/usr/bin/env python
 import sys, os
 try:
    import urllib.request as compat_urllib_request
 except ImportError: # Python 2
    import urllib2 as compat_urllib_request
 sys.stderr.write(u'Hi! We changed distribution method and now youtube-dl needs to update itself one more time.\n')
 sys.stderr.write(u'This will only happen once. Simply press enter to go on. Sorry for the trouble!\n')
 sys.stderr.write(u'The new location of the binaries is https://github.com/rg3/youtube-dl/downloads, not the git repository.\n\n')
 try:
 	raw_input()
 except NameError: # Python 3
 	input()
 filename = sys.argv[0]
 API_URL = "https://api.github.com/repos/rg3/youtube-dl/downloads"
 BIN_URL = "https://github.com/downloads/rg3/youtube-dl/youtube-dl"
 if not os.access(filename, os.W_OK):
    sys.exit('ERROR: no write permissions on %s' % filename)
 try:
    urlh = compat_urllib_request.urlopen(BIN_URL)
    newcontent = urlh.read()
    urlh.close()
 except (IOError, OSError) as err:
    sys.exit('ERROR: unable to download latest version')
 try:
    with open(filename, 'wb') as outf:
        outf.write(newcontent)
 except (IOError, OSError) as err:
    sys.exit('ERROR: unable to overwrite current version')
 sys.stderr.write(u'Done! Now you can run youtube-dl.\n')
--- a/devscripts/transition_helper_exe/setup.py
+++ b/devscripts/transition_helper_exe/setup.py
@ -1,12 +0,0 @@
 from distutils.core import setup
 import py2exe
 py2exe_options = {
    "bundle_files": 1,
    "compressed": 1,
    "optimize": 2,
    "dist_dir": '.',
    "dll_excludes": ['w9xpopen.exe']
 }
 setup(console=['youtube-dl.py'], options={ "py2exe": py2exe_options }, zipfile=None)
--- a/devscripts/transition_helper_exe/youtube-dl.py
+++ b/devscripts/transition_helper_exe/youtube-dl.py
@ -1,102 +0,0 @@
 #!/usr/bin/env python
 import sys, os
 import urllib2
 import json, hashlib
 def rsa_verify(message, signature, key):
    from struct import pack
    from hashlib import sha256
    from sys import version_info
    def b(x):
        if version_info[0] == 2: return x
        else: return x.encode('latin1')
    assert(type(message) == type(b('')))
    block_size = 0
    n = key[0]
    while n:
        block_size += 1
        n >>= 8
    signature = pow(int(signature, 16), key[1], key[0])
    raw_bytes = []
    while signature:
        raw_bytes.insert(0, pack("B", signature & 0xFF))
        signature >>= 8
    signature = (block_size - len(raw_bytes)) * b('\x00') + b('').join(raw_bytes)
    if signature[0:2] != b('\x00\x01'): return False
    signature = signature[2:]
    if not b('\x00') in signature: return False
    signature = signature[signature.index(b('\x00'))+1:]
    if not signature.startswith(b('\x30\x31\x30\x0D\x06\x09\x60\x86\x48\x01\x65\x03\x04\x02\x01\x05\x00\x04\x20')): return False
    signature = signature[19:]
    if signature != sha256(message).digest(): return False
    return True
 sys.stderr.write(u'Hi! We changed distribution method and now youtube-dl needs to update itself one more time.\n')
 sys.stderr.write(u'This will only happen once. Simply press enter to go on. Sorry for the trouble!\n')
 sys.stderr.write(u'From now on, get the binaries from http://rg3.github.com/youtube-dl/download.html, not from the git repository.\n\n')
 raw_input()
 filename = sys.argv[0]
 UPDATE_URL = "http://rg3.github.io/youtube-dl/update/"
 VERSION_URL = UPDATE_URL + 'LATEST_VERSION'
 JSON_URL = UPDATE_URL + 'versions.json'
 UPDATES_RSA_KEY = (0x9d60ee4d8f805312fdb15a62f87b95bd66177b91df176765d13514a0f1754bcd2057295c5b6f1d35daa6742c3ffc9a82d3e118861c207995a8031e151d863c9927e304576bc80692bc8e094896fcf11b66f3e29e04e3a71e9a11558558acea1840aec37fc396fb6b65dc81a1c4144e03bd1c011de62e3f1357b327d08426fe93, 65537)
 if not os.access(filename, os.W_OK):
    sys.exit('ERROR: no write permissions on %s' % filename)
 exe = os.path.abspath(filename)
 directory = os.path.dirname(exe)
 if not os.access(directory, os.W_OK):
    sys.exit('ERROR: no write permissions on %s' % directory)
 try:
    versions_info = urllib2.urlopen(JSON_URL).read().decode('utf-8')
    versions_info = json.loads(versions_info)
 except:
    sys.exit(u'ERROR: can\'t obtain versions info. Please try again later.')
 if not 'signature' in versions_info:
    sys.exit(u'ERROR: the versions file is not signed or corrupted. Aborting.')
 signature = versions_info['signature']
 del versions_info['signature']
 if not rsa_verify(json.dumps(versions_info, sort_keys=True), signature, UPDATES_RSA_KEY):
    sys.exit(u'ERROR: the versions file signature is invalid. Aborting.')
 version = versions_info['versions'][versions_info['latest']]
 try:
    urlh = urllib2.urlopen(version['exe'][0])
    newcontent = urlh.read()
    urlh.close()
 except (IOError, OSError) as err:
    sys.exit('ERROR: unable to download latest version')
 newcontent_hash = hashlib.sha256(newcontent).hexdigest()
 if newcontent_hash != version['exe'][1]:
    sys.exit(u'ERROR: the downloaded file hash does not match. Aborting.')
 try:
    with open(exe + '.new', 'wb') as outf:
        outf.write(newcontent)
 except (IOError, OSError) as err:
    sys.exit(u'ERROR: unable to write the new version')
 try:
    bat = os.path.join(directory, 'youtube-dl-updater.bat')
    b = open(bat, 'w')
    b.write("""
 echo Updating youtube-dl...
 ping 127.0.0.1 -n 5 -w 1000 > NUL
 move /Y "%s.new" "%s"
 del "%s"
    \n""" %(exe, exe, bat))
    b.close()
    os.startfile(bat)
 except (IOError, OSError) as err:
    sys.exit('ERROR: unable to overwrite current version')
 sys.stderr.write(u'Done! Now you can run youtube-dl.\n')
--- a/devscripts/zsh-completion.py
+++ b/devscripts/zsh-completion.py
@ -1,4 +1,6 @@
 #!/usr/bin/env python
 from __future__ import unicode_literals
 import os
 from os.path import dirname as dirn
 import sys
--- a/setup.py
+++ b/setup.py
@ -4,7 +4,6 @@
 from __future__ import print_function
 import os.path
 import pkg_resources
 import warnings
 import sys
@ -103,7 +102,9 @@ setup(
        "Programming Language :: Python :: 2.6",
        "Programming Language :: Python :: 2.7",
        "Programming Language :: Python :: 3",
-        "Programming Language :: Python :: 3.3"
+        "Programming Language :: Python :: 3.2",
        "Programming Language :: Python :: 3.3",
        "Programming Language :: Python :: 3.4",
    ],
    **params
--- a/test/helper.py
+++ b/test/helper.py
@ -57,7 +57,7 @@ class FakeYDL(YoutubeDL):
        # Different instances of the downloader can't share the same dictionary
        # some test set the "sublang" parameter, which would break the md5 checks.
        params = get_params(override=override)
-        super(FakeYDL, self).__init__(params)
+        super(FakeYDL, self).__init__(params, auto_init=False)
        self.result = []
    def to_screen(self, s, skip_eol=None):
@ -72,8 +72,10 @@ class FakeYDL(YoutubeDL):
    def expect_warning(self, regex):
        # Silence an expected warning matching a regex
        old_report_warning = self.report_warning
        def report_warning(self, message):
-            if re.match(regex, message): return
+            if re.match(regex, message):
                return
            old_report_warning(message)
        self.report_warning = types.MethodType(report_warning, self)
@ -145,7 +147,8 @@ def expect_info_dict(self, expected_dict, got_dict):
        info_dict_str = ''.join(
            '    %s: %s,\n' % (_repr(k), _repr(v))
            for k, v in test_info_dict.items())
-        write_string('\n"info_dict": {\n' + info_dict_str + '}\n', out=sys.stderr)
+        write_string(
            '\n\'info_dict\': {\n' + info_dict_str + '}\n', out=sys.stderr)
        self.assertFalse(
            missing_keys,
            'Missing keys in test definition: %s' % (
--- a/test/swftests/ConstArrayAccess.as
+++ b/test/swftests/ConstArrayAccess.as
@ -0,0 +1,18 @@
 // input: []
 // output: 4
 package {
 public class ConstArrayAccess {
 	private static const x:int = 2;
 	private static const ar:Array = ["42", "3411"];
    public static function main():int{
        var c:ConstArrayAccess = new ConstArrayAccess();
        return c.f();
    }
    public function f(): int {
    	return ar[1].length;
    }
 }
 }
--- a/test/swftests/ConstantInt.as
+++ b/test/swftests/ConstantInt.as
@ -0,0 +1,12 @@
 // input: []
 // output: 2
 package {
 public class ConstantInt {
 	private static const x:int = 2;
    public static function main():int{
        return x;
    }
 }
 }
--- a/test/swftests/DictCall.as
+++ b/test/swftests/DictCall.as
@ -0,0 +1,10 @@
 // input: [{"x": 1, "y": 2}]
 // output: 3
 package {
 public class DictCall {
    public static function main(d:Object):int{
        return d.x + d.y;
    }
 }
 }
--- a/test/swftests/EqualsOperator.as
+++ b/test/swftests/EqualsOperator.as
@ -0,0 +1,10 @@
 // input: []
 // output: false
 package {
 public class EqualsOperator {
    public static function main():Boolean{
        return 1 == 2;
    }
 }
 }
--- a/test/swftests/MemberAssignment.as
+++ b/test/swftests/MemberAssignment.as
@ -0,0 +1,22 @@
 // input: [1]
 // output: 2
 package {
 public class MemberAssignment {
    public var v:int;
    public function g():int {
        return this.v;
    }
    public function f(a:int):int{
        this.v = a;
        return this.v + this.g();
    }
    public static function main(a:int): int {
        var v:MemberAssignment = new MemberAssignment();
        return v.f(a);
    }
 }
 }
--- a/test/swftests/NeOperator.as
+++ b/test/swftests/NeOperator.as
@ -0,0 +1,24 @@
 // input: []
 // output: 123
 package {
 public class NeOperator {
    public static function main(): int {
        var res:int = 0;
        if (1 != 2) {
            res += 3;
        } else {
            res += 4;
        }
        if (2 != 2) {
            res += 10;
        } else {
            res += 20;
        }
        if (9 == 9) {
            res += 100;
        }
        return res;
    }
 }
 }
--- a/test/swftests/PrivateVoidCall.as
+++ b/test/swftests/PrivateVoidCall.as
@ -0,0 +1,22 @@
 // input: []
 // output: 9
 package {
 public class PrivateVoidCall {
    public static function main():int{
        var f:OtherClass = new OtherClass();
        f.func();
        return 9;
    }
 }
 }
 class OtherClass {
    private function pf():void {
        ;
    }
    public function func():void {
        this.pf();
    }
 }
--- a/test/swftests/StringBasics.as
+++ b/test/swftests/StringBasics.as
@ -0,0 +1,11 @@
 // input: []
 // output: 3
 package {
 public class StringBasics {
    public static function main():int{
        var s:String = "abc";
        return s.length;
    }
 }
 }
--- a/test/swftests/StringCharCodeAt.as
+++ b/test/swftests/StringCharCodeAt.as
@ -0,0 +1,11 @@
 // input: []
 // output: 9897
 package {
 public class StringCharCodeAt {
    public static function main():int{
        var s:String = "abc";
        return s.charCodeAt(1) * 100 + s.charCodeAt();
    }
 }
 }
--- a/test/swftests/StringConversion.as
+++ b/test/swftests/StringConversion.as
@ -0,0 +1,11 @@
 // input: []
 // output: 2
 package {
 public class StringConversion {
    public static function main():int{
        var s:String = String(99);
        return s.length;
    }
 }
 }
--- a/test/test_YoutubeDL.py
+++ b/test/test_YoutubeDL.py
@ -266,6 +266,7 @@ class TestFormatSelection(unittest.TestCase):
            'ext': 'mp4',
            'width': None,
        }
        def fname(templ):
            ydl = YoutubeDL({'outtmpl': templ})
            return ydl.prepare_filename(info)
--- a/test/test_age_restriction.py
+++ b/test/test_age_restriction.py
@ -1,4 +1,5 @@
 #!/usr/bin/env python
 from __future__ import unicode_literals
 # Allow direct execution
 import os
@ -19,7 +20,7 @@ def _download_restricted(url, filename, age):
        'age_limit': age,
        'skip_download': True,
        'writeinfojson': True,
-        "outtmpl": "%(id)s.%(ext)s",
+        'outtmpl': '%(id)s.%(ext)s',
    }
    ydl = YoutubeDL(params)
    ydl.add_default_info_extractors()
--- a/test/test_compat.py
+++ b/test/test_compat.py
@ -0,0 +1,46 @@
 #!/usr/bin/env python
 # coding: utf-8
 from __future__ import unicode_literals
 # Allow direct execution
 import os
 import sys
 import unittest
 sys.path.insert(0, os.path.dirname(os.path.dirname(os.path.abspath(__file__))))
 from youtube_dl.utils import get_filesystem_encoding
 from youtube_dl.compat import (
    compat_getenv,
    compat_expanduser,
 )
 class TestCompat(unittest.TestCase):
    def test_compat_getenv(self):
        test_str = 'тест'
        os.environ['YOUTUBE-DL-TEST'] = (
            test_str if sys.version_info >= (3, 0)
            else test_str.encode(get_filesystem_encoding()))
        self.assertEqual(compat_getenv('YOUTUBE-DL-TEST'), test_str)
    def test_compat_expanduser(self):
        old_home = os.environ.get('HOME')
        test_str = 'C:\Documents and Settings\тест\Application Data'
        os.environ['HOME'] = (
            test_str if sys.version_info >= (3, 0)
            else test_str.encode(get_filesystem_encoding()))
        self.assertEqual(compat_expanduser('~'), test_str)
        os.environ['HOME'] = old_home
    def test_all_present(self):
        import youtube_dl.compat
        all_names = youtube_dl.compat.__all__
        present_names = set(filter(
            lambda c: '_' in c and not c.startswith('_'),
            dir(youtube_dl.compat))) - set(['unicode_literals'])
        self.assertEqual(all_names, sorted(present_names))
 if __name__ == '__main__':
    unittest.main()
--- a/test/test_download.py
+++ b/test/test_download.py
@ -1,5 +1,7 @@
 #!/usr/bin/env python
 from __future__ import unicode_literals
 # Allow direct execution
 import os
 import sys
@ -23,10 +25,12 @@ import json
 import socket
 import youtube_dl.YoutubeDL
-from youtube_dl.utils import (
+from youtube_dl.compat import (
    compat_http_client,
    compat_urllib_error,
    compat_HTTPError,
 )
 from youtube_dl.utils import (
    DownloadError,
    ExtractorError,
    format_bytes,
@ -36,18 +40,22 @@ from youtube_dl.extractor import get_info_extractor
 RETRIES = 3
 class YoutubeDL(youtube_dl.YoutubeDL):
    def __init__(self, *args, **kwargs):
        self.to_stderr = self.to_screen
        self.processed_info_dicts = []
        super(YoutubeDL, self).__init__(*args, **kwargs)
    def report_warning(self, message):
        # Don't accept warnings during tests
        raise ExtractorError(message)
    def process_info(self, info_dict):
        self.processed_info_dicts.append(info_dict)
        return super(YoutubeDL, self).process_info(info_dict)
 def _file_md5(fn):
    with open(fn, 'rb') as f:
        return hashlib.md5(f.read()).hexdigest()
@ -57,10 +65,13 @@ defs = gettestcases()
 class TestDownload(unittest.TestCase):
    maxDiff = None
    def setUp(self):
        self.defs = defs
-### Dynamically generate tests
+# Dynamically generate tests
 def generator(test_case):
    def test_template(self):
@ -86,7 +97,7 @@ def generator(test_case):
            return
        for other_ie in other_ies:
            if not other_ie.working():
-                print_skipping(u'test depends on %sIE, marked as not WORKING' % other_ie.ie_key())
+                print_skipping('test depends on %sIE, marked as not WORKING' % other_ie.ie_key())
                return
        params = get_params(test_case.get('params', {}))
@ -94,9 +105,10 @@ def generator(test_case):
            params.setdefault('extract_flat', True)
            params.setdefault('skip_download', True)
-        ydl = YoutubeDL(params)
+        ydl = YoutubeDL(params, auto_init=False)
        ydl.add_default_info_extractors()
        finished_hook_called = set()
        def _hook(status):
            if status['status'] == 'finished':
                finished_hook_called.add(status['filename'])
@ -107,6 +119,7 @@ def generator(test_case):
            return tc.get('file') or ydl.prepare_filename(tc.get('info_dict', {}))
        res_dict = None
        def try_rm_tcs_files(tcs=None):
            if tcs is None:
                tcs = test_cases
@ -130,7 +143,7 @@ def generator(test_case):
                        raise
                    if try_num == RETRIES:
-                        report_warning(u'Failed due to network errors, skipping...')
+                        report_warning('Failed due to network errors, skipping...')
                        return
                    print('Retrying: {0} failed tries\n\n##########\n\n'.format(try_num))
@ -202,15 +215,15 @@ def generator(test_case):
    return test_template
-### And add them to TestDownload
+# And add them to TestDownload
 for n, test_case in enumerate(defs):
    test_method = generator(test_case)
    tname = 'test_' + str(test_case['name'])
    i = 1
    while hasattr(TestDownload, tname):
-        tname = 'test_'  + str(test_case['name']) + '_' + str(i)
+        tname = 'test_%s_%d' % (test_case['name'], i)
        i += 1
-    test_method.__name__ = tname
+    test_method.__name__ = str(tname)
    setattr(TestDownload, test_method.__name__, test_method)
    del test_method
--- a/test/test_execution.py
+++ b/test/test_execution.py
@ -1,3 +1,6 @@
 #!/usr/bin/env python
 from __future__ import unicode_literals
 import unittest
 import sys
@ -6,11 +9,13 @@ import subprocess
 rootDir = os.path.dirname(os.path.dirname(os.path.abspath(__file__)))
 try:
    _DEV_NULL = subprocess.DEVNULL
 except AttributeError:
    _DEV_NULL = open(os.devnull, 'wb')
 class TestExecution(unittest.TestCase):
    def test_import(self):
        subprocess.check_call([sys.executable, '-c', 'import youtube_dl'], cwd=rootDir)
--- a/test/test_subtitles.py
+++ b/test/test_subtitles.py
@ -1,4 +1,5 @@
 #!/usr/bin/env python
 from __future__ import unicode_literals
 # Allow direct execution
 import os
@ -22,6 +23,7 @@ from youtube_dl.extractor import (
 class BaseTestSubtitles(unittest.TestCase):
    url = None
    IE = None
    def setUp(self):
        self.DL = FakeYDL()
        self.ie = self.IE(self.DL)
@ -74,7 +76,7 @@ class TestYoutubeSubtitles(BaseTestSubtitles):
        self.assertEqual(md5(subtitles['en']), '3cb210999d3e021bd6c7f0ea751eab06')
    def test_youtube_list_subtitles(self):
-        self.DL.expect_warning(u'Video doesn\'t have automatic captions')
+        self.DL.expect_warning('Video doesn\'t have automatic captions')
        self.DL.params['listsubtitles'] = True
        info_dict = self.getInfoDict()
        self.assertEqual(info_dict, None)
@ -87,7 +89,7 @@ class TestYoutubeSubtitles(BaseTestSubtitles):
        self.assertTrue(subtitles['it'] is not None)
    def test_youtube_nosubtitles(self):
-        self.DL.expect_warning(u'video doesn\'t have subtitles')
+        self.DL.expect_warning('video doesn\'t have subtitles')
        self.url = 'n5BB19UTcdA'
        self.DL.params['writesubtitles'] = True
        self.DL.params['allsubtitles'] = True
@ -101,7 +103,7 @@ class TestYoutubeSubtitles(BaseTestSubtitles):
        self.DL.params['subtitleslangs'] = langs
        subtitles = self.getSubtitles()
        for lang in langs:
-            self.assertTrue(subtitles.get(lang) is not None, u'Subtitles for \'%s\' not extracted' % lang)
+            self.assertTrue(subtitles.get(lang) is not None, 'Subtitles for \'%s\' not extracted' % lang)
 class TestDailymotionSubtitles(BaseTestSubtitles):
@ -130,20 +132,20 @@ class TestDailymotionSubtitles(BaseTestSubtitles):
        self.assertEqual(len(subtitles.keys()), 5)
    def test_list_subtitles(self):
-        self.DL.expect_warning(u'Automatic Captions not supported by this server')
+        self.DL.expect_warning('Automatic Captions not supported by this server')
        self.DL.params['listsubtitles'] = True
        info_dict = self.getInfoDict()
        self.assertEqual(info_dict, None)
    def test_automatic_captions(self):
-        self.DL.expect_warning(u'Automatic Captions not supported by this server')
+        self.DL.expect_warning('Automatic Captions not supported by this server')
        self.DL.params['writeautomaticsub'] = True
        self.DL.params['subtitleslang'] = ['en']
        subtitles = self.getSubtitles()
        self.assertTrue(len(subtitles.keys()) == 0)
    def test_nosubtitles(self):
-        self.DL.expect_warning(u'video doesn\'t have subtitles')
+        self.DL.expect_warning('video doesn\'t have subtitles')
        self.url = 'http://www.dailymotion.com/video/x12u166_le-zapping-tele-star-du-08-aout-2013_tv'
        self.DL.params['writesubtitles'] = True
        self.DL.params['allsubtitles'] = True
@ -156,7 +158,7 @@ class TestDailymotionSubtitles(BaseTestSubtitles):
        self.DL.params['subtitleslangs'] = langs
        subtitles = self.getSubtitles()
        for lang in langs:
-            self.assertTrue(subtitles.get(lang) is not None, u'Subtitles for \'%s\' not extracted' % lang)
+            self.assertTrue(subtitles.get(lang) is not None, 'Subtitles for \'%s\' not extracted' % lang)
 class TestTedSubtitles(BaseTestSubtitles):
@ -185,13 +187,13 @@ class TestTedSubtitles(BaseTestSubtitles):
        self.assertTrue(len(subtitles.keys()) >= 28)
    def test_list_subtitles(self):
-        self.DL.expect_warning(u'Automatic Captions not supported by this server')
+        self.DL.expect_warning('Automatic Captions not supported by this server')
        self.DL.params['listsubtitles'] = True
        info_dict = self.getInfoDict()
        self.assertEqual(info_dict, None)
    def test_automatic_captions(self):
-        self.DL.expect_warning(u'Automatic Captions not supported by this server')
+        self.DL.expect_warning('Automatic Captions not supported by this server')
        self.DL.params['writeautomaticsub'] = True
        self.DL.params['subtitleslang'] = ['en']
        subtitles = self.getSubtitles()
@ -203,7 +205,7 @@ class TestTedSubtitles(BaseTestSubtitles):
        self.DL.params['subtitleslangs'] = langs
        subtitles = self.getSubtitles()
        for lang in langs:
-            self.assertTrue(subtitles.get(lang) is not None, u'Subtitles for \'%s\' not extracted' % lang)
+            self.assertTrue(subtitles.get(lang) is not None, 'Subtitles for \'%s\' not extracted' % lang)
 class TestBlipTVSubtitles(BaseTestSubtitles):
@ -211,13 +213,13 @@ class TestBlipTVSubtitles(BaseTestSubtitles):
    IE = BlipTVIE
    def test_list_subtitles(self):
-        self.DL.expect_warning(u'Automatic Captions not supported by this server')
+        self.DL.expect_warning('Automatic Captions not supported by this server')
        self.DL.params['listsubtitles'] = True
        info_dict = self.getInfoDict()
        self.assertEqual(info_dict, None)
    def test_allsubtitles(self):
-        self.DL.expect_warning(u'Automatic Captions not supported by this server')
+        self.DL.expect_warning('Automatic Captions not supported by this server')
        self.DL.params['writesubtitles'] = True
        self.DL.params['allsubtitles'] = True
        subtitles = self.getSubtitles()
@ -251,20 +253,20 @@ class TestVimeoSubtitles(BaseTestSubtitles):
        self.assertEqual(set(subtitles.keys()), set(['de', 'en', 'es', 'fr']))
    def test_list_subtitles(self):
-        self.DL.expect_warning(u'Automatic Captions not supported by this server')
+        self.DL.expect_warning('Automatic Captions not supported by this server')
        self.DL.params['listsubtitles'] = True
        info_dict = self.getInfoDict()
        self.assertEqual(info_dict, None)
    def test_automatic_captions(self):
-        self.DL.expect_warning(u'Automatic Captions not supported by this server')
+        self.DL.expect_warning('Automatic Captions not supported by this server')
        self.DL.params['writeautomaticsub'] = True
        self.DL.params['subtitleslang'] = ['en']
        subtitles = self.getSubtitles()
        self.assertTrue(len(subtitles.keys()) == 0)
    def test_nosubtitles(self):
-        self.DL.expect_warning(u'video doesn\'t have subtitles')
+        self.DL.expect_warning('video doesn\'t have subtitles')
        self.url = 'http://vimeo.com/56015672'
        self.DL.params['writesubtitles'] = True
        self.DL.params['allsubtitles'] = True
@ -277,7 +279,7 @@ class TestVimeoSubtitles(BaseTestSubtitles):
        self.DL.params['subtitleslangs'] = langs
        subtitles = self.getSubtitles()
        for lang in langs:
-            self.assertTrue(subtitles.get(lang) is not None, u'Subtitles for \'%s\' not extracted' % lang)
+            self.assertTrue(subtitles.get(lang) is not None, 'Subtitles for \'%s\' not extracted' % lang)
 class TestWallaSubtitles(BaseTestSubtitles):
@ -285,13 +287,13 @@ class TestWallaSubtitles(BaseTestSubtitles):
    IE = WallaIE
    def test_list_subtitles(self):
-        self.DL.expect_warning(u'Automatic Captions not supported by this server')
+        self.DL.expect_warning('Automatic Captions not supported by this server')
        self.DL.params['listsubtitles'] = True
        info_dict = self.getInfoDict()
        self.assertEqual(info_dict, None)
    def test_allsubtitles(self):
-        self.DL.expect_warning(u'Automatic Captions not supported by this server')
+        self.DL.expect_warning('Automatic Captions not supported by this server')
        self.DL.params['writesubtitles'] = True
        self.DL.params['allsubtitles'] = True
        subtitles = self.getSubtitles()
@ -299,7 +301,7 @@ class TestWallaSubtitles(BaseTestSubtitles):
        self.assertEqual(md5(subtitles['heb']), 'e758c5d7cb982f6bef14f377ec7a3920')
    def test_nosubtitles(self):
-        self.DL.expect_warning(u'video doesn\'t have subtitles')
+        self.DL.expect_warning('video doesn\'t have subtitles')
        self.url = 'http://vod.walla.co.il/movie/2642630/one-direction-all-for-one'
        self.DL.params['writesubtitles'] = True
        self.DL.params['allsubtitles'] = True
--- a/test/test_swfinterp.py
+++ b/test/test_swfinterp.py
@ -1,4 +1,5 @@
 #!/usr/bin/env python
 from __future__ import unicode_literals
 # Allow direct execution
 import os
@ -37,7 +38,9 @@ def _make_testfunc(testfile):
                or os.path.getmtime(swf_file) < os.path.getmtime(as_file)):
            # Recompile
            try:
-                subprocess.check_call(['mxmlc', '-output', swf_file, as_file])
+                subprocess.check_call([
                    'mxmlc', '-output', swf_file,
                    '-static-link-runtime-shared-libraries', as_file])
            except OSError as ose:
                if ose.errno == errno.ENOENT:
                    print('mxmlc not found! Skipping test.')
--- a/test/test_unicode_literals.py
+++ b/test/test_unicode_literals.py
@ -9,14 +9,13 @@ rootDir = os.path.dirname(os.path.dirname(os.path.abspath(__file__)))
 IGNORED_FILES = [
    'setup.py',  # http://bugs.python.org/issue13943
    'conf.py',
    'buildserver.py',
 ]
 class TestUnicodeLiterals(unittest.TestCase):
    def test_all_files(self):
        print('Skipping this test (not yet fully implemented)')
        return
        for dirpath, _, filenames in os.walk(rootDir):
            for basename in filenames:
                if not basename.endswith('.py'):
@ -30,10 +29,10 @@ class TestUnicodeLiterals(unittest.TestCase):
                if "'" not in code and '"' not in code:
                    continue
-                imps = 'from __future__ import unicode_literals'
+                self.assertRegexpMatches(
-                self.assertTrue(
+                    code,
-                    imps in code,
+                    r'(?:#.*\n*)?from __future__ import (?:[a-z_]+,\s*)*unicode_literals',
-                    ' %s  missing in %s' % (imps, fn))
+                    'unicode_literals import  missing in %s' % fn)
                m = re.search(r'(?<=\s)u[\'"](?!\)|,|$)', code)
                if m is not None:
--- a/test/test_utils.py
+++ b/test/test_utils.py
@ -16,11 +16,11 @@ import json
 import xml.etree.ElementTree
 from youtube_dl.utils import (
    clean_html,
    DateRange,
    encodeFilename,
    find_xpath_attr,
    fix_xml_ampersands,
    get_meta_content,
    orderedSet,
    OnDemandPagedList,
    InAdvancePagedList,
@ -45,9 +45,9 @@ from youtube_dl.utils import (
    escape_rfc3986,
    escape_url,
    js_to_json,
-    get_filesystem_encoding,
+    intlist_to_bytes,
-    compat_getenv,
+    args_to_str,
-    compat_expanduser,
+    parse_filesize,
 )
@ -157,17 +157,6 @@ class TestUtil(unittest.TestCase):
        self.assertEqual(find_xpath_attr(doc, './/node', 'x', 'a'), doc[1])
        self.assertEqual(find_xpath_attr(doc, './/node', 'y', 'c'), doc[2])
    def test_meta_parser(self):
        testhtml = '''
        <head>
            <meta name="description" content="foo &amp; bar">
            <meta content='Plato' name='author'/>
        </head>
        '''
        get_meta = lambda name: get_meta_content(name, testhtml)
        self.assertEqual(get_meta('description'), 'foo & bar')
        self.assertEqual(get_meta('author'), 'Plato')
    def test_xpath_with_ns(self):
        testxml = '''<root xmlns:media="http://example.com/">
            <media:song>
@ -182,7 +171,7 @@ class TestUtil(unittest.TestCase):
        self.assertEqual(find('media:song/url').text, 'http://server.com/download.mp3')
    def test_smuggle_url(self):
-        data = {u"ö": u"ö", u"abc": [3]}
+        data = {"ö": "ö", "abc": [3]}
        url = 'https://foo.bar/baz?x=y#a'
        smug_url = smuggle_url(url, data)
        unsmug_url, unsmug_data = unsmuggle_url(smug_url)
@ -230,6 +219,7 @@ class TestUtil(unittest.TestCase):
        self.assertEqual(parse_duration('0m0s'), 0)
        self.assertEqual(parse_duration('0s'), 0)
        self.assertEqual(parse_duration('01:02:03.05'), 3723.05)
        self.assertEqual(parse_duration('T30M38S'), 1838)
    def test_fix_xml_ampersands(self):
        self.assertEqual(
@ -296,6 +286,10 @@ class TestUtil(unittest.TestCase):
        d = json.loads(stripped)
        self.assertEqual(d, [{"id": "532cb", "x": 3}])
        stripped = strip_jsonp('parseMetadata({"STATUS":"OK"})\n\n\n//epc')
        d = json.loads(stripped)
        self.assertEqual(d, {'STATUS': 'OK'})
    def test_uppercase_escape(self):
        self.assertEqual(uppercase_escape('aä'), 'aä')
        self.assertEqual(uppercase_escape('\\U0001d550'), '𝕐')
@ -359,17 +353,29 @@ class TestUtil(unittest.TestCase):
        on = js_to_json('{"abc": true}')
        self.assertEqual(json.loads(on), {'abc': True})
-    def test_compat_getenv(self):
+    def test_clean_html(self):
-        test_str = 'тест'
+        self.assertEqual(clean_html('a:\nb'), 'a: b')
-        os.environ['YOUTUBE-DL-TEST'] = (test_str if sys.version_info >= (3, 0)
+        self.assertEqual(clean_html('a:\n   "b"'), 'a:    "b"')
            else test_str.encode(get_filesystem_encoding()))
        self.assertEqual(compat_getenv('YOUTUBE-DL-TEST'), test_str)
-    def test_compat_expanduser(self):
+    def test_intlist_to_bytes(self):
-        test_str = 'C:\Documents and Settings\тест\Application Data'
+        self.assertEqual(
-        os.environ['HOME'] = (test_str if sys.version_info >= (3, 0)
+            intlist_to_bytes([0, 1, 127, 128, 255]),
-            else test_str.encode(get_filesystem_encoding()))
+            b'\x00\x01\x7f\x80\xff')
-        self.assertEqual(compat_expanduser('~'), test_str)
+
    def test_args_to_str(self):
        self.assertEqual(
            args_to_str(['foo', 'ba/r', '-baz', '2 be', '']),
            'foo ba/r -baz \'2 be\' \'\''
        )
    def test_parse_filesize(self):
        self.assertEqual(parse_filesize(None), None)
        self.assertEqual(parse_filesize(''), None)
        self.assertEqual(parse_filesize('91 B'), 91)
        self.assertEqual(parse_filesize('foobar'), None)
        self.assertEqual(parse_filesize('2 MiB'), 2097152)
        self.assertEqual(parse_filesize('5 GB'), 5000000000)
        self.assertEqual(parse_filesize('1.2Tb'), 1200000000000)
 if __name__ == '__main__':
    unittest.main()
--- a/test/test_write_annotations.py
+++ b/test/test_write_annotations.py
@ -1,5 +1,6 @@
 #!/usr/bin/env python
 # coding: utf-8
 from __future__ import unicode_literals
 # Allow direct execution
 import os
@ -31,17 +32,16 @@ params = get_params({
 })
 TEST_ID = 'gr51aVj-mLg'
 ANNOTATIONS_FILE = TEST_ID + '.flv.annotations.xml'
 EXPECTED_ANNOTATIONS = ['Speech bubble', 'Note', 'Title', 'Spotlight', 'Label']
 class TestAnnotations(unittest.TestCase):
    def setUp(self):
        # Clear old files
        self.tearDown()
    def test_info_json(self):
        expected = list(EXPECTED_ANNOTATIONS)  # Two annotations could have the same text.
        ie = youtube_dl.extractor.YoutubeIE()
@ -71,7 +71,6 @@ class TestAnnotations(unittest.TestCase):
        # We should have seen (and removed) all the expected annotation texts.
        self.assertEqual(len(expected), 0, 'Not all expected annotations were found.')
    def tearDown(self):
        try_rm(ANNOTATIONS_FILE)
--- a/test/test_write_info_json.py
+++ b/test/test_write_info_json.py
@ -1,5 +1,6 @@
 #!/usr/bin/env python
 # coding: utf-8
 from __future__ import unicode_literals
 # Allow direct execution
 import os
@ -32,7 +33,7 @@ params = get_params({
 TEST_ID = 'BaW_jenozKc'
 INFO_JSON_FILE = TEST_ID + '.info.json'
 DESCRIPTION_FILE = TEST_ID + '.mp4.description'
-EXPECTED_DESCRIPTION = u'''test chars:  "'/\ä↭𝕐
+EXPECTED_DESCRIPTION = '''test chars:  "'/\ä↭𝕐
 test URL: https://github.com/rg3/youtube-dl/issues/1892
 This is a test video for youtube-dl.
@ -53,11 +54,11 @@ class TestInfoJSON(unittest.TestCase):
        self.assertTrue(os.path.exists(INFO_JSON_FILE))
        with io.open(INFO_JSON_FILE, 'r', encoding='utf-8') as jsonf:
            jd = json.load(jsonf)
-        self.assertEqual(jd['upload_date'], u'20121002')
+        self.assertEqual(jd['upload_date'], '20121002')
        self.assertEqual(jd['description'], EXPECTED_DESCRIPTION)
        self.assertEqual(jd['id'], TEST_ID)
        self.assertEqual(jd['extractor'], 'youtube')
-        self.assertEqual(jd['title'], u'''youtube-dl test video "'/\ä↭𝕐''')
+        self.assertEqual(jd['title'], '''youtube-dl test video "'/\ä↭𝕐''')
        self.assertEqual(jd['uploader'], 'Philipp Hagemeister')
        self.assertTrue(os.path.exists(DESCRIPTION_FILE))
--- a/test/test_youtube_lists.py
+++ b/test/test_youtube_lists.py
@ -1,4 +1,5 @@
 #!/usr/bin/env python
 from __future__ import unicode_literals
 # Allow direct execution
 import os
@ -12,10 +13,6 @@ from test.helper import FakeYDL
 from youtube_dl.extractor import (
    YoutubePlaylistIE,
    YoutubeIE,
    YoutubeChannelIE,
    YoutubeShowIE,
    YoutubeTopListIE,
    YoutubeSearchURLIE,
 )
--- a/test/test_youtube_signature.py
+++ b/test/test_youtube_signature.py
@ -14,7 +14,7 @@ import re
 import string
 from youtube_dl.extractor import YoutubeIE
-from youtube_dl.utils import compat_str, compat_urlretrieve
+from youtube_dl.compat import compat_str, compat_urlretrieve
 _TESTS = [
    (
--- a/youtube_dl/YoutubeDL.py
+++ b/youtube_dl/YoutubeDL.py
@ -22,13 +22,15 @@ import traceback
 if os.name == 'nt':
    import ctypes
-from .utils import (
+from .compat import (
    compat_cookiejar,
    compat_expanduser,
    compat_http_client,
    compat_str,
    compat_urllib_error,
    compat_urllib_request,
 )
 from .utils import (
    escape_url,
    ContentTooShortError,
    date_from_str,
@ -58,10 +60,12 @@ from .utils import (
    write_string,
    YoutubeDLHandler,
    prepend_extension,
    args_to_str,
 )
 from .cache import Cache
 from .extractor import get_info_extractor, gen_extractors
 from .downloader import get_suitable_downloader
 from .downloader.rtmp import rtmpdump_version
 from .postprocessor import FFmpegMergerPP, FFmpegPostProcessor
 from .version import __version__
@ -250,6 +254,22 @@ class YoutubeDL(object):
            self.print_debug_header()
            self.add_default_info_extractors()
    def warn_if_short_id(self, argv):
        # short YouTube ID starting with dash?
        idxs = [
            i for i, a in enumerate(argv)
            if re.match(r'^-[0-9A-Za-z_-]{10}$', a)]
        if idxs:
            correct_argv = (
                ['youtube-dl'] +
                [a for i, a in enumerate(argv) if i not in idxs] +
                ['--'] + [argv[i] for i in idxs]
            )
            self.report_warning(
                'Long argument string detected. '
                'Use -- to separate parameters and URLs, like this:\n%s\n' %
                args_to_str(correct_argv))
    def add_info_extractor(self, ie):
        """Add an InfoExtractor object to the end of the list."""
        self._ies.append(ie)
@ -621,7 +641,7 @@ class YoutubeDL(object):
            return self.process_ie_result(
                new_result, download=download, extra_info=extra_info)
-        elif result_type == 'playlist':
+        elif result_type == 'playlist' or result_type == 'multi_video':
            # We process each entry in the playlist
            playlist = ie_result.get('title', None) or ie_result.get('id', None)
            self.to_screen('[download] Downloading playlist: %s' % playlist)
@ -655,6 +675,8 @@ class YoutubeDL(object):
                extra = {
                    'n_entries': n_entries,
                    'playlist': playlist,
                    'playlist_id': ie_result.get('id'),
                    'playlist_title': ie_result.get('title'),
                    'playlist_index': i + playliststart,
                    'extractor': ie_result['extractor'],
                    'webpage_url': ie_result['webpage_url'],
@ -674,14 +696,20 @@ class YoutubeDL(object):
            ie_result['entries'] = playlist_results
            return ie_result
        elif result_type == 'compat_list':
            self.report_warning(
                'Extractor %s returned a compat_list result. '
                'It needs to be updated.' % ie_result.get('extractor'))
            def _fixup(r):
-                self.add_extra_info(r,
+                self.add_extra_info(
                    r,
                    {
                        'extractor': ie_result['extractor'],
                        'webpage_url': ie_result['webpage_url'],
                        'webpage_url_basename': url_basename(ie_result['webpage_url']),
                        'extractor_key': ie_result['extractor_key'],
-                    })
+                    }
                )
                return r
            ie_result['entries'] = [
                self.process_ie_result(_fixup(r), download, extra_info)
@ -833,6 +861,13 @@ class YoutubeDL(object):
                        formats_info = (self.select_format(format_1, formats),
                                        self.select_format(format_2, formats))
                        if all(formats_info):
                            # The first format must contain the video and the
                            # second the audio
                            if formats_info[0].get('vcodec') == 'none':
                                self.report_error('The first format must '
                                                  'contain the video, try using '
                                                  '"-f %s+%s"' % (format_2, format_1))
                                return
                            selected_format = {
                                'requested_formats': formats_info,
                                'format': rf,
@ -989,7 +1024,7 @@ class YoutubeDL(object):
            else:
                self.to_screen('[info] Writing video description metadata as JSON to: ' + infofn)
                try:
-                    write_json_file(info_dict, encodeFilename(infofn))
+                    write_json_file(info_dict, infofn)
                except (OSError, IOError):
                    self.report_error('Cannot write metadata to JSON file ' + infofn)
                    return
@ -1294,11 +1329,13 @@ class YoutubeDL(object):
            self.report_warning(
                'Your Python is broken! Update to a newer and supported version')
        stdout_encoding = getattr(
            sys.stdout, 'encoding', 'missing (%s)' % type(sys.stdout).__name__)
        encoding_str = (
            '[debug] Encodings: locale %s, fs %s, out %s, pref %s\n' % (
                locale.getpreferredencoding(),
                sys.getfilesystemencoding(),
-                sys.stdout.encoding,
+                stdout_encoding,
                self.get_encoding()))
        write_string(encoding_str, encoding=None)
@ -1321,6 +1358,7 @@ class YoutubeDL(object):
            platform.python_version(), platform_name()))
        exe_versions = FFmpegPostProcessor.get_versions()
        exe_versions['rtmpdump'] = rtmpdump_version()
        exe_str = ', '.join(
            '%s %s' % (exe, v)
            for exe, v in sorted(exe_versions.items())
--- a/youtube_dl/init.py
+++ b/youtube_dl/init.py
@ -1,6 +1,8 @@
 #!/usr/bin/env python
 # -*- coding: utf-8 -*-
 from __future__ import unicode_literals
 __license__ = 'Public Domain'
 import codecs
@ -13,10 +15,13 @@ import sys
 from .options import (
    parseOpts,
 )
-from .utils import (
+from .compat import (
    compat_expanduser,
    compat_getpass,
    compat_print,
    workaround_optparse_bug9161,
 )
 from .utils import (
    DateRange,
    DEFAULT_OUTTMPL,
    decodeOption,
@ -53,7 +58,9 @@ def _real_main(argv=None):
        # https://github.com/rg3/youtube-dl/issues/820
        codecs.register(lambda name: codecs.lookup('utf-8') if name == 'cp65001' else None)
-    setproctitle(u'youtube-dl')
+    workaround_optparse_bug9161()
    setproctitle('youtube-dl')
    parser, opts, args = parseOpts(argv)
@ -69,10 +76,10 @@ def _real_main(argv=None):
    if opts.headers is not None:
        for h in opts.headers:
            if h.find(':', 1) < 0:
-                parser.error(u'wrong header formatting, it should be key:value, not "%s"'%h)
+                parser.error('wrong header formatting, it should be key:value, not "%s"' % h)
            key, value = h.split(':', 2)
            if opts.verbose:
-                write_string(u'[debug] Adding header from command line option %s:%s\n'%(key, value))
+                write_string('[debug] Adding header from command line option %s:%s\n' % (key, value))
            std_headers[key] = value
    # Dump user agent
@ -90,9 +97,9 @@ def _real_main(argv=None):
                batchfd = io.open(opts.batchfile, 'r', encoding='utf-8', errors='ignore')
            batch_urls = read_batch_urls(batchfd)
            if opts.verbose:
-                write_string(u'[debug] Batch file urls: ' + repr(batch_urls) + u'\n')
+                write_string('[debug] Batch file urls: ' + repr(batch_urls) + '\n')
        except IOError:
-            sys.exit(u'ERROR: batch file could not be read')
+            sys.exit('ERROR: batch file could not be read')
    all_urls = batch_urls + args
    all_urls = [url.strip() for url in all_urls]
    _enc = preferredencoding()
@ -105,7 +112,7 @@ def _real_main(argv=None):
            compat_print(ie.IE_NAME + (' (CURRENTLY BROKEN)' if not ie._WORKING else ''))
            matchedUrls = [url for url in all_urls if ie.suitable(url)]
            for mu in matchedUrls:
-                compat_print(u'  ' + mu)
+                compat_print('  ' + mu)
        sys.exit(0)
    if opts.list_extractor_descriptions:
        for ie in sorted(extractors, key=lambda ie: ie.IE_NAME.lower()):
@ -115,63 +122,62 @@ def _real_main(argv=None):
            if desc is False:
                continue
            if hasattr(ie, 'SEARCH_KEY'):
-                _SEARCHES = (u'cute kittens', u'slithering pythons', u'falling cat', u'angry poodle', u'purple fish', u'running tortoise', u'sleeping bunny')
+                _SEARCHES = ('cute kittens', 'slithering pythons', 'falling cat', 'angry poodle', 'purple fish', 'running tortoise', 'sleeping bunny')
-                _COUNTS = (u'', u'5', u'10', u'all')
+                _COUNTS = ('', '5', '10', 'all')
-                desc += u' (Example: "%s%s:%s" )' % (ie.SEARCH_KEY, random.choice(_COUNTS), random.choice(_SEARCHES))
+                desc += ' (Example: "%s%s:%s" )' % (ie.SEARCH_KEY, random.choice(_COUNTS), random.choice(_SEARCHES))
            compat_print(desc)
        sys.exit(0)
    # Conflicting, missing and erroneous options
    if opts.usenetrc and (opts.username is not None or opts.password is not None):
-        parser.error(u'using .netrc conflicts with giving username/password')
+        parser.error('using .netrc conflicts with giving username/password')
    if opts.password is not None and opts.username is None:
-        parser.error(u'account username missing\n')
+        parser.error('account username missing\n')
    if opts.outtmpl is not None and (opts.usetitle or opts.autonumber or opts.useid):
-        parser.error(u'using output template conflicts with using title, video ID or auto number')
+        parser.error('using output template conflicts with using title, video ID or auto number')
    if opts.usetitle and opts.useid:
-        parser.error(u'using title conflicts with using video ID')
+        parser.error('using title conflicts with using video ID')
    if opts.username is not None and opts.password is None:
-        opts.password = compat_getpass(u'Type account password and press [Return]: ')
+        opts.password = compat_getpass('Type account password and press [Return]: ')
    if opts.ratelimit is not None:
        numeric_limit = FileDownloader.parse_bytes(opts.ratelimit)
        if numeric_limit is None:
-            parser.error(u'invalid rate limit specified')
+            parser.error('invalid rate limit specified')
        opts.ratelimit = numeric_limit
    if opts.min_filesize is not None:
        numeric_limit = FileDownloader.parse_bytes(opts.min_filesize)
        if numeric_limit is None:
-            parser.error(u'invalid min_filesize specified')
+            parser.error('invalid min_filesize specified')
        opts.min_filesize = numeric_limit
    if opts.max_filesize is not None:
        numeric_limit = FileDownloader.parse_bytes(opts.max_filesize)
        if numeric_limit is None:
-            parser.error(u'invalid max_filesize specified')
+            parser.error('invalid max_filesize specified')
        opts.max_filesize = numeric_limit
    if opts.retries is not None:
        try:
            opts.retries = int(opts.retries)
        except (TypeError, ValueError):
-            parser.error(u'invalid retry count specified')
+            parser.error('invalid retry count specified')
    if opts.buffersize is not None:
        numeric_buffersize = FileDownloader.parse_bytes(opts.buffersize)
        if numeric_buffersize is None:
-            parser.error(u'invalid buffer size specified')
+            parser.error('invalid buffer size specified')
        opts.buffersize = numeric_buffersize
    if opts.playliststart <= 0:
-        raise ValueError(u'Playlist start must be positive')
+        raise ValueError('Playlist start must be positive')
    if opts.playlistend not in (-1, None) and opts.playlistend < opts.playliststart:
-        raise ValueError(u'Playlist end must be greater than playlist start')
+        raise ValueError('Playlist end must be greater than playlist start')
    if opts.extractaudio:
        if opts.audioformat not in ['best', 'aac', 'mp3', 'm4a', 'opus', 'vorbis', 'wav']:
-            parser.error(u'invalid audio format specified')
+            parser.error('invalid audio format specified')
    if opts.audioquality:
        opts.audioquality = opts.audioquality.strip('k').strip('K')
        if not opts.audioquality.isdigit():
-            parser.error(u'invalid audio quality specified')
+            parser.error('invalid audio quality specified')
    if opts.recodevideo is not None:
        if opts.recodevideo not in ['mp4', 'flv', 'webm', 'ogg', 'mkv']:
-            parser.error(u'invalid video recode format specified')
+            parser.error('invalid video recode format specified')
    if opts.date is not None:
        date = DateRange.day(opts.date)
    else:
@ -183,7 +189,7 @@ def _real_main(argv=None):
    # --all-sub automatically sets --write-sub if --write-auto-sub is not given
    # this was the old behaviour if only --all-sub was given.
-    if opts.allsubtitles and (opts.writeautomaticsub == False):
+    if opts.allsubtitles and not opts.writeautomaticsub:
        opts.writesubtitles = True
    if sys.version_info < (3,):
@ -191,17 +197,17 @@ def _real_main(argv=None):
        if opts.outtmpl is not None:
            opts.outtmpl = opts.outtmpl.decode(preferredencoding())
    outtmpl = ((opts.outtmpl is not None and opts.outtmpl)
-            or (opts.format == '-1' and opts.usetitle and u'%(title)s-%(id)s-%(format)s.%(ext)s')
+               or (opts.format == '-1' and opts.usetitle and '%(title)s-%(id)s-%(format)s.%(ext)s')
-            or (opts.format == '-1' and u'%(id)s-%(format)s.%(ext)s')
+               or (opts.format == '-1' and '%(id)s-%(format)s.%(ext)s')
-            or (opts.usetitle and opts.autonumber and u'%(autonumber)s-%(title)s-%(id)s.%(ext)s')
+               or (opts.usetitle and opts.autonumber and '%(autonumber)s-%(title)s-%(id)s.%(ext)s')
-            or (opts.usetitle and u'%(title)s-%(id)s.%(ext)s')
+               or (opts.usetitle and '%(title)s-%(id)s.%(ext)s')
-            or (opts.useid and u'%(id)s.%(ext)s')
+               or (opts.useid and '%(id)s.%(ext)s')
-            or (opts.autonumber and u'%(autonumber)s-%(id)s.%(ext)s')
+               or (opts.autonumber and '%(autonumber)s-%(id)s.%(ext)s')
               or DEFAULT_OUTTMPL)
    if not os.path.splitext(outtmpl)[1] and opts.extractaudio:
-        parser.error(u'Cannot download a video and extract audio into the same'
+        parser.error('Cannot download a video and extract audio into the same'
-                     u' file! Use "{0}.%(ext)s" instead of "{0}" as the output'
+                     ' file! Use "{0}.%(ext)s" instead of "{0}" as the output'
-                     u' template'.format(outtmpl))
+                     ' template'.format(outtmpl))
    any_printing = opts.geturl or opts.gettitle or opts.getid or opts.getthumbnail or opts.getdescription or opts.getfilename or opts.getformat or opts.getduration or opts.dumpjson or opts.dump_single_json
    download_archive_fn = compat_expanduser(opts.download_archive) if opts.download_archive is not None else opts.download_archive
@ -310,7 +316,6 @@ def _real_main(argv=None):
                ydl.add_post_processor(FFmpegAudioFixPP())
            ydl.add_post_processor(AtomicParsleyPP())
        # Please keep ExecAfterDownload towards the bottom as it allows the user to modify the final file in any way.
        # So if the user is able to remove the file before your postprocessor runs it might cause a few problems.
        if opts.exec_cmd:
@ -327,18 +332,19 @@ def _real_main(argv=None):
        # Maybe do nothing
        if (len(all_urls) < 1) and (opts.load_info_filename is None):
-            if not (opts.update_self or opts.rm_cachedir):
+            if opts.update_self or opts.rm_cachedir:
                parser.error(u'you must provide at least one URL')
            else:
                sys.exit()
            ydl.warn_if_short_id(sys.argv[1:] if argv is None else argv)
            parser.error('you must provide at least one URL')
        try:
            if opts.load_info_filename is not None:
                retcode = ydl.download_with_info_file(opts.load_info_filename)
            else:
                retcode = ydl.download(all_urls)
        except MaxDownloadsReached:
-            ydl.to_screen(u'--max-download limit reached, aborting.')
+            ydl.to_screen('--max-download limit reached, aborting.')
            retcode = 101
    sys.exit(retcode)
@ -350,6 +356,6 @@ def main(argv=None):
    except DownloadError:
        sys.exit(1)
    except SameFileError:
-        sys.exit(u'ERROR: fixed output name but more than one file to download')
+        sys.exit('ERROR: fixed output name but more than one file to download')
    except KeyboardInterrupt:
-        sys.exit(u'\nERROR: Interrupted by user')
+        sys.exit('\nERROR: Interrupted by user')
--- a/youtube_dl/main.py
+++ b/youtube_dl/main.py
@ -1,4 +1,5 @@
 #!/usr/bin/env python
 from __future__ import unicode_literals
 # Execute with
 # $ python youtube_dl/__main__.py (2.6+)
--- a/youtube_dl/aes.py
+++ b/youtube_dl/aes.py
@ -1,3 +1,5 @@
 from __future__ import unicode_literals
 __all__ = ['aes_encrypt', 'key_expansion', 'aes_ctr_decrypt', 'aes_cbc_decrypt', 'aes_decrypt_text']
 import base64
@ -7,6 +9,7 @@ from .utils import bytes_to_intlist, intlist_to_bytes
 BLOCK_SIZE_BYTES = 16
 def aes_ctr_decrypt(data, key, counter):
    """
    Decrypt with aes in counter mode
@ -32,6 +35,7 @@ def aes_ctr_decrypt(data, key, counter):
    return decrypted_data
 def aes_cbc_decrypt(data, key, iv):
    """
    Decrypt with aes in CBC mode
@ -57,6 +61,7 @@ def aes_cbc_decrypt(data, key, iv):
    return decrypted_data
 def key_expansion(data):
    """
    Generate key schedule
@ -91,6 +96,7 @@ def key_expansion(data):
    return data
 def aes_encrypt(data, expanded_key):
    """
    Encrypt one block with aes
@ -111,6 +117,7 @@ def aes_encrypt(data, expanded_key):
    return data
 def aes_decrypt(data, expanded_key):
    """
    Decrypt one block with aes
@ -131,6 +138,7 @@ def aes_decrypt(data, expanded_key):
    return data
 def aes_decrypt_text(data, password, key_size_bytes):
    """
    Decrypt text
@ -157,6 +165,7 @@ def aes_decrypt_text(data, password, key_size_bytes):
    class Counter:
        __value = nonce + [0] * (BLOCK_SIZE_BYTES - NONCE_LENGTH_BYTES)
        def next_value(self):
            temp = self.__value
            self.__value = inc(self.__value)
@ -241,15 +250,19 @@ RIJNDAEL_LOG_TABLE = (0x00, 0x00, 0x19, 0x01, 0x32, 0x02, 0x1a, 0xc6, 0x4b, 0xc7
                      0x44, 0x11, 0x92, 0xd9, 0x23, 0x20, 0x2e, 0x89, 0xb4, 0x7c, 0xb8, 0x26, 0x77, 0x99, 0xe3, 0xa5,
                      0x67, 0x4a, 0xed, 0xde, 0xc5, 0x31, 0xfe, 0x18, 0x0d, 0x63, 0x8c, 0x80, 0xc0, 0xf7, 0x70, 0x07)
 def sub_bytes(data):
    return [SBOX[x] for x in data]
 def sub_bytes_inv(data):
    return [SBOX_INV[x] for x in data]
 def rotate(data):
    return data[1:] + [data[0]]
 def key_schedule_core(data, rcon_iteration):
    data = rotate(data)
    data = sub_bytes(data)
@ -257,14 +270,17 @@ def key_schedule_core(data, rcon_iteration):
    return data
 def xor(data1, data2):
    return [x ^ y for x, y in zip(data1, data2)]
 def rijndael_mul(a, b):
    if(a == 0 or b == 0):
        return 0
    return RIJNDAEL_EXP_TABLE[(RIJNDAEL_LOG_TABLE[a] + RIJNDAEL_LOG_TABLE[b]) % 0xFF]
 def mix_column(data, matrix):
    data_mixed = []
    for row in range(4):
@ -275,6 +291,7 @@ def mix_column(data, matrix):
        data_mixed.append(mixed)
    return data_mixed
 def mix_columns(data, matrix=MIX_COLUMN_MATRIX):
    data_mixed = []
    for i in range(4):
@ -282,9 +299,11 @@ def mix_columns(data, matrix=MIX_COLUMN_MATRIX):
        data_mixed += mix_column(column, matrix)
    return data_mixed
 def mix_columns_inv(data):
    return mix_columns(data, MIX_COLUMN_MATRIX_INV)
 def shift_rows(data):
    data_shifted = []
    for column in range(4):
@ -292,6 +311,7 @@ def shift_rows(data):
            data_shifted.append(data[((column + row) & 0b11) * 4 + row])
    return data_shifted
 def shift_rows_inv(data):
    data_shifted = []
    for column in range(4):
@ -299,6 +319,7 @@ def shift_rows_inv(data):
            data_shifted.append(data[((column - row) & 0b11) * 4 + row])
    return data_shifted
 def inc(data):
    data = data[:]  # copy
    for i in range(len(data) - 1, -1, -1):
--- a/youtube_dl/cache.py
+++ b/youtube_dl/cache.py
@ -8,10 +8,8 @@ import re
 import shutil
 import traceback
-from .utils import (
+from .compat import compat_expanduser, compat_getenv
-    compat_expanduser,
+from .utils import write_json_file
    write_json_file,
 )
 class Cache(object):
@ -21,7 +19,7 @@ class Cache(object):
    def _get_root_dir(self):
        res = self._ydl.params.get('cachedir')
        if res is None:
-            cache_root = os.environ.get('XDG_CACHE_HOME', '~/.cache')
+            cache_root = compat_getenv('XDG_CACHE_HOME', '~/.cache')
            res = os.path.join(cache_root, 'youtube-dl')
        return compat_expanduser(res)
--- a/youtube_dl/compat.py
+++ b/youtube_dl/compat.py
@ -0,0 +1,356 @@
 from __future__ import unicode_literals
 import getpass
 import optparse
 import os
 import re
 import subprocess
 import sys
 try:
    import urllib.request as compat_urllib_request
 except ImportError:  # Python 2
    import urllib2 as compat_urllib_request
 try:
    import urllib.error as compat_urllib_error
 except ImportError:  # Python 2
    import urllib2 as compat_urllib_error
 try:
    import urllib.parse as compat_urllib_parse
 except ImportError:  # Python 2
    import urllib as compat_urllib_parse
 try:
    from urllib.parse import urlparse as compat_urllib_parse_urlparse
 except ImportError:  # Python 2
    from urlparse import urlparse as compat_urllib_parse_urlparse
 try:
    import urllib.parse as compat_urlparse
 except ImportError:  # Python 2
    import urlparse as compat_urlparse
 try:
    import http.cookiejar as compat_cookiejar
 except ImportError:  # Python 2
    import cookielib as compat_cookiejar
 try:
    import html.entities as compat_html_entities
 except ImportError:  # Python 2
    import htmlentitydefs as compat_html_entities
 try:
    import html.parser as compat_html_parser
 except ImportError:  # Python 2
    import HTMLParser as compat_html_parser
 try:
    import http.client as compat_http_client
 except ImportError:  # Python 2
    import httplib as compat_http_client
 try:
    from urllib.error import HTTPError as compat_HTTPError
 except ImportError:  # Python 2
    from urllib2 import HTTPError as compat_HTTPError
 try:
    from urllib.request import urlretrieve as compat_urlretrieve
 except ImportError:  # Python 2
    from urllib import urlretrieve as compat_urlretrieve
 try:
    from subprocess import DEVNULL
    compat_subprocess_get_DEVNULL = lambda: DEVNULL
 except ImportError:
    compat_subprocess_get_DEVNULL = lambda: open(os.path.devnull, 'w')
 try:
    from urllib.parse import unquote as compat_urllib_parse_unquote
 except ImportError:
    def compat_urllib_parse_unquote(string, encoding='utf-8', errors='replace'):
        if string == '':
            return string
        res = string.split('%')
        if len(res) == 1:
            return string
        if encoding is None:
            encoding = 'utf-8'
        if errors is None:
            errors = 'replace'
        # pct_sequence: contiguous sequence of percent-encoded bytes, decoded
        pct_sequence = b''
        string = res[0]
        for item in res[1:]:
            try:
                if not item:
                    raise ValueError
                pct_sequence += item[:2].decode('hex')
                rest = item[2:]
                if not rest:
                    # This segment was just a single percent-encoded character.
                    # May be part of a sequence of code units, so delay decoding.
                    # (Stored in pct_sequence).
                    continue
            except ValueError:
                rest = '%' + item
            # Encountered non-percent-encoded characters. Flush the current
            # pct_sequence.
            string += pct_sequence.decode(encoding, errors) + rest
            pct_sequence = b''
        if pct_sequence:
            # Flush the final pct_sequence
            string += pct_sequence.decode(encoding, errors)
        return string
 try:
    from urllib.parse import parse_qs as compat_parse_qs
 except ImportError:  # Python 2
    # HACK: The following is the correct parse_qs implementation from cpython 3's stdlib.
    # Python 2's version is apparently totally broken
    def _parse_qsl(qs, keep_blank_values=False, strict_parsing=False,
                   encoding='utf-8', errors='replace'):
        qs, _coerce_result = qs, unicode
        pairs = [s2 for s1 in qs.split('&') for s2 in s1.split(';')]
        r = []
        for name_value in pairs:
            if not name_value and not strict_parsing:
                continue
            nv = name_value.split('=', 1)
            if len(nv) != 2:
                if strict_parsing:
                    raise ValueError("bad query field: %r" % (name_value,))
                # Handle case of a control-name with no equal sign
                if keep_blank_values:
                    nv.append('')
                else:
                    continue
            if len(nv[1]) or keep_blank_values:
                name = nv[0].replace('+', ' ')
                name = compat_urllib_parse_unquote(
                    name, encoding=encoding, errors=errors)
                name = _coerce_result(name)
                value = nv[1].replace('+', ' ')
                value = compat_urllib_parse_unquote(
                    value, encoding=encoding, errors=errors)
                value = _coerce_result(value)
                r.append((name, value))
        return r
    def compat_parse_qs(qs, keep_blank_values=False, strict_parsing=False,
                        encoding='utf-8', errors='replace'):
        parsed_result = {}
        pairs = _parse_qsl(qs, keep_blank_values, strict_parsing,
                           encoding=encoding, errors=errors)
        for name, value in pairs:
            if name in parsed_result:
                parsed_result[name].append(value)
            else:
                parsed_result[name] = [value]
        return parsed_result
 try:
    compat_str = unicode  # Python 2
 except NameError:
    compat_str = str
 try:
    compat_chr = unichr  # Python 2
 except NameError:
    compat_chr = chr
 try:
    from xml.etree.ElementTree import ParseError as compat_xml_parse_error
 except ImportError:  # Python 2.6
    from xml.parsers.expat import ExpatError as compat_xml_parse_error
 try:
    from shlex import quote as shlex_quote
 except ImportError:  # Python < 3.3
    def shlex_quote(s):
        if re.match(r'^[-_\w./]+$', s):
            return s
        else:
            return "'" + s.replace("'", "'\"'\"'") + "'"
 def compat_ord(c):
    if type(c) is int:
        return c
    else:
        return ord(c)
 if sys.version_info >= (3, 0):
    compat_getenv = os.getenv
    compat_expanduser = os.path.expanduser
 else:
    # Environment variables should be decoded with filesystem encoding.
    # Otherwise it will fail if any non-ASCII characters present (see #3854 #3217 #2918)
    def compat_getenv(key, default=None):
        from .utils import get_filesystem_encoding
        env = os.getenv(key, default)
        if env:
            env = env.decode(get_filesystem_encoding())
        return env
    # HACK: The default implementations of os.path.expanduser from cpython do not decode
    # environment variables with filesystem encoding. We will work around this by
    # providing adjusted implementations.
    # The following are os.path.expanduser implementations from cpython 2.7.8 stdlib
    # for different platforms with correct environment variables decoding.
    if os.name == 'posix':
        def compat_expanduser(path):
            """Expand ~ and ~user constructions.  If user or $HOME is unknown,
            do nothing."""
            if not path.startswith('~'):
                return path
            i = path.find('/', 1)
            if i < 0:
                i = len(path)
            if i == 1:
                if 'HOME' not in os.environ:
                    import pwd
                    userhome = pwd.getpwuid(os.getuid()).pw_dir
                else:
                    userhome = compat_getenv('HOME')
            else:
                import pwd
                try:
                    pwent = pwd.getpwnam(path[1:i])
                except KeyError:
                    return path
                userhome = pwent.pw_dir
            userhome = userhome.rstrip('/')
            return (userhome + path[i:]) or '/'
    elif os.name == 'nt' or os.name == 'ce':
        def compat_expanduser(path):
            """Expand ~ and ~user constructs.
            If user or $HOME is unknown, do nothing."""
            if path[:1] != '~':
                return path
            i, n = 1, len(path)
            while i < n and path[i] not in '/\\':
                i = i + 1
            if 'HOME' in os.environ:
                userhome = compat_getenv('HOME')
            elif 'USERPROFILE' in os.environ:
                userhome = compat_getenv('USERPROFILE')
            elif not 'HOMEPATH' in os.environ:
                return path
            else:
                try:
                    drive = compat_getenv('HOMEDRIVE')
                except KeyError:
                    drive = ''
                userhome = os.path.join(drive, compat_getenv('HOMEPATH'))
            if i != 1:  # ~user
                userhome = os.path.join(os.path.dirname(userhome), path[1:i])
            return userhome + path[i:]
    else:
        compat_expanduser = os.path.expanduser
 if sys.version_info < (3, 0):
    def compat_print(s):
        from .utils import preferredencoding
        print(s.encode(preferredencoding(), 'xmlcharrefreplace'))
 else:
    def compat_print(s):
        assert isinstance(s, compat_str)
        print(s)
 try:
    subprocess_check_output = subprocess.check_output
 except AttributeError:
    def subprocess_check_output(*args, **kwargs):
        assert 'input' not in kwargs
        p = subprocess.Popen(*args, stdout=subprocess.PIPE, **kwargs)
        output, _ = p.communicate()
        ret = p.poll()
        if ret:
            raise subprocess.CalledProcessError(ret, p.args, output=output)
        return output
 if sys.version_info < (3, 0) and sys.platform == 'win32':
    def compat_getpass(prompt, *args, **kwargs):
        if isinstance(prompt, compat_str):
            from .utils import preferredencoding
            prompt = prompt.encode(preferredencoding())
        return getpass.getpass(prompt, *args, **kwargs)
 else:
    compat_getpass = getpass.getpass
 # Old 2.6 and 2.7 releases require kwargs to be bytes
 try:
    (lambda x: x)(**{'x': 0})
 except TypeError:
    def compat_kwargs(kwargs):
        return dict((bytes(k), v) for k, v in kwargs.items())
 else:
    compat_kwargs = lambda kwargs: kwargs
 # Fix https://github.com/rg3/youtube-dl/issues/4223
 # See http://bugs.python.org/issue9161 for what is broken
 def workaround_optparse_bug9161():
    op = optparse.OptionParser()
    og = optparse.OptionGroup(op, 'foo')
    try:
        og.add_option('-t')
    except TypeError:
        real_add_option = optparse.OptionGroup.add_option
        def _compat_add_option(self, *args, **kwargs):
            enc = lambda v: (
                v.encode('ascii', 'replace') if isinstance(v, compat_str)
                else v)
            bargs = [enc(a) for a in args]
            bkwargs = dict(
                (k, enc(v)) for k, v in kwargs.items())
            return real_add_option(self, *bargs, **bkwargs)
        optparse.OptionGroup.add_option = _compat_add_option
 __all__ = [
    'compat_HTTPError',
    'compat_chr',
    'compat_cookiejar',
    'compat_expanduser',
    'compat_getenv',
    'compat_getpass',
    'compat_html_entities',
    'compat_html_parser',
    'compat_http_client',
    'compat_kwargs',
    'compat_ord',
    'compat_parse_qs',
    'compat_print',
    'compat_str',
    'compat_subprocess_get_DEVNULL',
    'compat_urllib_error',
    'compat_urllib_parse',
    'compat_urllib_parse_unquote',
    'compat_urllib_parse_urlparse',
    'compat_urllib_request',
    'compat_urlparse',
    'compat_urlretrieve',
    'compat_xml_parse_error',
    'shlex_quote',
    'subprocess_check_output',
    'workaround_optparse_bug9161',
 ]
--- a/youtube_dl/downloader/init.py
+++ b/youtube_dl/downloader/init.py
@ -30,3 +30,8 @@ def get_suitable_downloader(info_dict):
        return F4mFD
    else:
        return HttpFD
 __all__ = [
    'get_suitable_downloader',
    'FileDownloader',
 ]
--- a/youtube_dl/downloader/common.py
+++ b/youtube_dl/downloader/common.py
@ -1,3 +1,5 @@
 from __future__ import unicode_literals
 import os
 import re
 import sys
@ -159,14 +161,14 @@ class FileDownloader(object):
    def temp_name(self, filename):
        """Returns a temporary filename for the given filename."""
-        if self.params.get('nopart', False) or filename == u'-' or \
+        if self.params.get('nopart', False) or filename == '-' or \
                (os.path.exists(encodeFilename(filename)) and not os.path.isfile(encodeFilename(filename))):
            return filename
-        return filename + u'.part'
+        return filename + '.part'
    def undo_temp_name(self, filename):
-        if filename.endswith(u'.part'):
+        if filename.endswith('.part'):
-            return filename[:-len(u'.part')]
+            return filename[:-len('.part')]
        return filename
    def try_rename(self, old_filename, new_filename):
@ -175,7 +177,7 @@ class FileDownloader(object):
                return
            os.rename(encodeFilename(old_filename), encodeFilename(new_filename))
        except (IOError, OSError) as err:
-            self.report_error(u'unable to rename file: %s' % compat_str(err))
+            self.report_error('unable to rename file: %s' % compat_str(err))
    def try_utime(self, filename, last_modified_hdr):
        """Try to set the last-modified time of the given file."""
@ -200,10 +202,10 @@ class FileDownloader(object):
    def report_destination(self, filename):
        """Report destination filename."""
-        self.to_screen(u'[download] Destination: ' + filename)
+        self.to_screen('[download] Destination: ' + filename)
    def _report_progress_status(self, msg, is_last_line=False):
-        fullmsg = u'[download] ' + msg
+        fullmsg = '[download] ' + msg
        if self.params.get('progress_with_newline', False):
            self.to_screen(fullmsg)
        else:
@ -211,13 +213,13 @@ class FileDownloader(object):
                prev_len = getattr(self, '_report_progress_prev_line_length',
                                   0)
                if prev_len > len(fullmsg):
-                    fullmsg += u' ' * (prev_len - len(fullmsg))
+                    fullmsg += ' ' * (prev_len - len(fullmsg))
                self._report_progress_prev_line_length = len(fullmsg)
-                clear_line = u'\r'
+                clear_line = '\r'
            else:
-                clear_line = (u'\r\x1b[K' if sys.stderr.isatty() else u'\r')
+                clear_line = ('\r\x1b[K' if sys.stderr.isatty() else '\r')
            self.to_screen(clear_line + fullmsg, skip_eol=not is_last_line)
-        self.to_console_title(u'youtube-dl ' + msg)
+        self.to_console_title('youtube-dl ' + msg)
    def report_progress(self, percent, data_len_str, speed, eta):
        """Report download progress."""
@ -233,7 +235,7 @@ class FileDownloader(object):
            percent_str = 'Unknown %'
        speed_str = self.format_speed(speed)
-        msg = (u'%s of %s at %s ETA %s' %
+        msg = ('%s of %s at %s ETA %s' %
               (percent_str, data_len_str, speed_str, eta_str))
        self._report_progress_status(msg)
@ -243,37 +245,37 @@ class FileDownloader(object):
        downloaded_str = format_bytes(downloaded_data_len)
        speed_str = self.format_speed(speed)
        elapsed_str = FileDownloader.format_seconds(elapsed)
-        msg = u'%s at %s (%s)' % (downloaded_str, speed_str, elapsed_str)
+        msg = '%s at %s (%s)' % (downloaded_str, speed_str, elapsed_str)
        self._report_progress_status(msg)
    def report_finish(self, data_len_str, tot_time):
        """Report download finished."""
        if self.params.get('noprogress', False):
-            self.to_screen(u'[download] Download completed')
+            self.to_screen('[download] Download completed')
        else:
            self._report_progress_status(
-                (u'100%% of %s in %s' %
+                ('100%% of %s in %s' %
                 (data_len_str, self.format_seconds(tot_time))),
                is_last_line=True)
    def report_resuming_byte(self, resume_len):
        """Report attempt to resume at given byte."""
-        self.to_screen(u'[download] Resuming download at byte %s' % resume_len)
+        self.to_screen('[download] Resuming download at byte %s' % resume_len)
    def report_retry(self, count, retries):
        """Report retry in case of HTTP error 5xx"""
-        self.to_screen(u'[download] Got server HTTP error. Retrying (attempt %d of %d)...' % (count, retries))
+        self.to_screen('[download] Got server HTTP error. Retrying (attempt %d of %d)...' % (count, retries))
    def report_file_already_downloaded(self, file_name):
        """Report file has already been fully downloaded."""
        try:
-            self.to_screen(u'[download] %s has already been downloaded' % file_name)
+            self.to_screen('[download] %s has already been downloaded' % file_name)
        except UnicodeEncodeError:
-            self.to_screen(u'[download] The file has already been downloaded')
+            self.to_screen('[download] The file has already been downloaded')
    def report_unable_to_resume(self):
        """Report it was impossible to resume download."""
-        self.to_screen(u'[download] Unable to resume')
+        self.to_screen('[download] Unable to resume')
    def download(self, filename, info_dict):
        """Download to a filename using the info from info_dict
@ -293,7 +295,7 @@ class FileDownloader(object):
    def real_download(self, filename, info_dict):
        """Real download process. Redefine in subclasses."""
-        raise NotImplementedError(u'This method must be implemented by subclasses')
+        raise NotImplementedError('This method must be implemented by subclasses')
    def _hook_progress(self, status):
        for ph in self._progress_hooks:
--- a/youtube_dl/downloader/f4m.py
+++ b/youtube_dl/downloader/f4m.py
@ -225,13 +225,15 @@ class F4mFD(FileDownloader):
        self.to_screen('[download] Downloading f4m manifest')
        manifest = self.ydl.urlopen(man_url).read()
        self.report_destination(filename)
-        http_dl = HttpQuietDownloader(self.ydl,
+        http_dl = HttpQuietDownloader(
            self.ydl,
            {
                'continuedl': True,
                'quiet': True,
                'noprogress': True,
                'test': self.params.get('test', False),
-            })
+            }
        )
        doc = etree.fromstring(manifest)
        formats = [(int(f.attrib.get('bitrate', -1)), f) for f in doc.findall(_add_ns('media'))]
--- a/youtube_dl/downloader/hls.py
+++ b/youtube_dl/downloader/hls.py
@ -28,14 +28,14 @@ class HlsFD(FileDownloader):
            if check_executable(program, ['-version']):
                break
        else:
-            self.report_error(u'm3u8 download detected but ffmpeg or avconv could not be found. Please install one.')
+            self.report_error('m3u8 download detected but ffmpeg or avconv could not be found. Please install one.')
            return False
        cmd = [program] + args
        retval = subprocess.call(cmd)
        if retval == 0:
            fsize = os.path.getsize(encodeFilename(tmpfilename))
-            self.to_screen(u'\r[%s] %s bytes' % (cmd[0], fsize))
+            self.to_screen('\r[%s] %s bytes' % (cmd[0], fsize))
            self.try_rename(tmpfilename, filename)
            self._hook_progress({
                'downloaded_bytes': fsize,
@ -45,8 +45,8 @@ class HlsFD(FileDownloader):
            })
            return True
        else:
-            self.to_stderr(u"\n")
+            self.to_stderr('\n')
-            self.report_error(u'%s exited with code %d' % (program, retval))
+            self.report_error('%s exited with code %d' % (program, retval))
            return False
@ -101,4 +101,3 @@ class NativeHlsFD(FileDownloader):
        })
        self.try_rename(tmpfilename, filename)
        return True
--- a/youtube_dl/downloader/http.py
+++ b/youtube_dl/downloader/http.py
@ -1,3 +1,5 @@
 from __future__ import unicode_literals
 import os
 import time
@ -106,7 +108,7 @@ class HttpFD(FileDownloader):
                self.report_retry(count, retries)
        if count > retries:
-            self.report_error(u'giving up after %s retries' % retries)
+            self.report_error('giving up after %s retries' % retries)
            return False
        data_len = data.info().get('Content-length', None)
@ -124,10 +126,10 @@ class HttpFD(FileDownloader):
            min_data_len = self.params.get("min_filesize", None)
            max_data_len = self.params.get("max_filesize", None)
            if min_data_len is not None and data_len < min_data_len:
-                self.to_screen(u'\r[download] File is smaller than min-filesize (%s bytes < %s bytes). Aborting.' % (data_len, min_data_len))
+                self.to_screen('\r[download] File is smaller than min-filesize (%s bytes < %s bytes). Aborting.' % (data_len, min_data_len))
                return False
            if max_data_len is not None and data_len > max_data_len:
-                self.to_screen(u'\r[download] File is larger than max-filesize (%s bytes > %s bytes). Aborting.' % (data_len, max_data_len))
+                self.to_screen('\r[download] File is larger than max-filesize (%s bytes > %s bytes). Aborting.' % (data_len, max_data_len))
                return False
        data_len_str = format_bytes(data_len)
@ -151,13 +153,13 @@ class HttpFD(FileDownloader):
                    filename = self.undo_temp_name(tmpfilename)
                    self.report_destination(filename)
                except (OSError, IOError) as err:
-                    self.report_error(u'unable to open for writing: %s' % str(err))
+                    self.report_error('unable to open for writing: %s' % str(err))
                    return False
            try:
                stream.write(data_block)
            except (IOError, OSError) as err:
-                self.to_stderr(u"\n")
+                self.to_stderr('\n')
-                self.report_error(u'unable to write data: %s' % str(err))
+                self.report_error('unable to write data: %s' % str(err))
                return False
            if not self.params.get('noresizebuffer', False):
                block_size = self.best_block_size(after - before, len(data_block))
@ -188,10 +190,10 @@ class HttpFD(FileDownloader):
            self.slow_down(start, byte_counter - resume_len)
        if stream is None:
-            self.to_stderr(u"\n")
+            self.to_stderr('\n')
-            self.report_error(u'Did not get any data blocks')
+            self.report_error('Did not get any data blocks')
            return False
-        if tmpfilename != u'-':
+        if tmpfilename != '-':
            stream.close()
        self.report_finish(data_len_str, (time.time() - start))
        if data_len is not None and byte_counter != data_len:
--- a/youtube_dl/downloader/mplayer.py
+++ b/youtube_dl/downloader/mplayer.py
@ -1,7 +1,10 @@
 from __future__ import unicode_literals
 import os
 import subprocess
 from .common import FileDownloader
 from ..compat import compat_subprocess_get_DEVNULL
 from ..utils import (
    encodeFilename,
 )
@ -13,19 +16,23 @@ class MplayerFD(FileDownloader):
        self.report_destination(filename)
        tmpfilename = self.temp_name(filename)
-        args = ['mplayer', '-really-quiet', '-vo', 'null', '-vc', 'dummy', '-dumpstream', '-dumpfile', tmpfilename, url]
+        args = [
            'mplayer', '-really-quiet', '-vo', 'null', '-vc', 'dummy',
            '-dumpstream', '-dumpfile', tmpfilename, url]
        # Check for mplayer first
        try:
-            subprocess.call(['mplayer', '-h'], stdout=(open(os.path.devnull, 'w')), stderr=subprocess.STDOUT)
+            subprocess.call(
                ['mplayer', '-h'],
                stdout=compat_subprocess_get_DEVNULL(), stderr=subprocess.STDOUT)
        except (OSError, IOError):
-            self.report_error(u'MMS or RTSP download detected but "%s" could not be run' % args[0])
+            self.report_error('MMS or RTSP download detected but "%s" could not be run' % args[0])
            return False
        # Download using mplayer.
        retval = subprocess.call(args)
        if retval == 0:
            fsize = os.path.getsize(encodeFilename(tmpfilename))
-            self.to_screen(u'\r[%s] %s bytes' % (args[0], fsize))
+            self.to_screen('\r[%s] %s bytes' % (args[0], fsize))
            self.try_rename(tmpfilename, filename)
            self._hook_progress({
                'downloaded_bytes': fsize,
@ -35,6 +42,6 @@ class MplayerFD(FileDownloader):
            })
            return True
        else:
-            self.to_stderr(u"\n")
+            self.to_stderr('\n')
-            self.report_error(u'mplayer exited with code %d' % retval)
+            self.report_error('mplayer exited with code %d' % retval)
            return False
--- a/youtube_dl/downloader/rtmp.py
+++ b/youtube_dl/downloader/rtmp.py
@ -12,9 +12,15 @@ from ..utils import (
    compat_str,
    encodeFilename,
    format_bytes,
    get_exe_version,
 )
 def rtmpdump_version():
    return get_exe_version(
        'rtmpdump', ['--help'], r'(?i)RTMPDump\s*v?([0-9a-zA-Z._-]+)')
 class RtmpFD(FileDownloader):
    def real_download(self, filename, info_dict):
        def run_rtmpdump(args):
--- a/youtube_dl/extractor/init.py
+++ b/youtube_dl/extractor/init.py
@ -1,3 +1,5 @@
 from __future__ import unicode_literals
 from .abc import ABCIE
 from .academicearth import AcademicEarthCourseIE
 from .addanime import AddAnimeIE
@ -32,6 +34,7 @@ from .bilibili import BiliBiliIE
 from .blinkx import BlinkxIE
 from .bliptv import BlipTVIE, BlipTVUserIE
 from .bloomberg import BloombergIE
 from .bpb import BpbIE
 from .br import BRIE
 from .breakcom import BreakIE
 from .brightcove import BrightcoveIE
@ -115,6 +118,7 @@ from .fktv import (
    FKTVPosteckeIE,
 )
 from .flickr import FlickrIE
 from .folketinget import FolketingetIE
 from .fourtube import FourTubeIE
 from .franceculture import FranceCultureIE
 from .franceinter import FranceInterIE
@ -127,6 +131,7 @@ from .francetv import (
 )
 from .freesound import FreesoundIE
 from .freespeech import FreespeechIE
 from .freevideo import FreeVideoIE
 from .funnyordie import FunnyOrDieIE
 from .gamekings import GamekingsIE
 from .gameone import (
@ -141,6 +146,7 @@ from .generic import GenericIE
 from .glide import GlideIE
 from .globo import GloboIE
 from .godtube import GodTubeIE
 from .goldenmoustache import GoldenMoustacheIE
 from .golem import GolemIE
 from .googleplus import GooglePlusIE
 from .googlesearch import GoogleSearchIE
@ -323,6 +329,7 @@ from .sbs import SBSIE
 from .scivee import SciVeeIE
 from .screencast import ScreencastIE
 from .servingsys import ServingSysIE
 from .sexu import SexuIE
 from .sexykarma import SexyKarmaIE
 from .shared import SharedIE
 from .sharesix import ShareSixIE
@ -368,6 +375,7 @@ from .syfy import SyfyIE
 from .sztvhu import SztvHuIE
 from .tagesschau import TagesschauIE
 from .tapely import TapelyIE
 from .tass import TassIE
 from .teachertube import (
    TeacherTubeIE,
    TeacherTubeUserIE,
@ -376,6 +384,7 @@ from .teachingchannel import TeachingChannelIE
 from .teamcoco import TeamcocoIE
 from .techtalks import TechTalksIE
 from .ted import TEDIE
 from .telebruxelles import TeleBruxellesIE
 from .telecinco import TelecincoIE
 from .telemb import TeleMBIE
 from .tenplay import TenPlayIE
@ -387,6 +396,7 @@ from .thesixtyone import TheSixtyOneIE
 from .thisav import ThisAVIE
 from .tinypic import TinyPicIE
 from .tlc import TlcIE, TlcDeIE
 from .tmz import TMZIE
 from .tnaflix import TNAFlixIE
 from .thvideo import (
    THVideoIE,
@ -400,6 +410,7 @@ from .trutube import TruTubeIE
 from .tube8 import Tube8IE
 from .tudou import TudouIE
 from .tumblr import TumblrIE
 from .tunein import TuneInIE
 from .turbo import TurboIE
 from .tutv import TutvIE
 from .tvigle import TvigleIE
@ -421,6 +432,7 @@ from .vesti import VestiIE
 from .vevo import VevoIE
 from .vgtv import VGTVIE
 from .vh1 import VH1IE
 from .vice import ViceIE
 from .viddler import ViddlerIE
 from .videobam import VideoBamIE
 from .videodetective import VideoDetectiveIE
@ -448,7 +460,10 @@ from .vine import (
    VineUserIE,
 )
 from .viki import VikiIE
-from .vk import VKIE
+from .vk import (
    VKIE,
    VKUserVideosIE,
 )
 from .vodlocker import VodlockerIE
 from .vporn import VpornIE
 from .vrt import VRTIE
@ -472,6 +487,7 @@ from .wrzuta import WrzutaIE
 from .xbef import XBefIE
 from .xboxclips import XboxClipsIE
 from .xhamster import XHamsterIE
 from .xminus import XMinusIE
 from .xnxx import XNXXIE
 from .xvideos import XVideosIE
 from .xtube import XTubeUserIE, XTubeIE
@ -502,6 +518,10 @@ from .youtube import (
    YoutubeWatchLaterIE,
 )
 from .zdf import ZDFIE
 from .zingmp3 import (
    ZingMp3SongIE,
    ZingMp3AlbumIE,
 )
 _ALL_CLASSES = [
    klass
--- a/youtube_dl/extractor/abc.py
+++ b/youtube_dl/extractor/abc.py
@ -11,13 +11,13 @@ class ABCIE(InfoExtractor):
    _VALID_URL = r'http://www\.abc\.net\.au/news/[^/]+/[^/]+/(?P<id>\d+)'
    _TEST = {
-        'url': 'http://www.abc.net.au/news/2014-07-25/bringing-asylum-seekers-to-australia-would-give/5624716',
+        'url': 'http://www.abc.net.au/news/2014-11-05/australia-to-staff-ebola-treatment-centre-in-sierra-leone/5868334',
-        'md5': 'dad6f8ad011a70d9ddf887ce6d5d0742',
+        'md5': 'cb3dd03b18455a661071ee1e28344d9f',
        'info_dict': {
-            'id': '5624716',
+            'id': '5868334',
            'ext': 'mp4',
-            'title': 'Bringing asylum seekers to Australia would give them right to asylum claims: professor',
+            'title': 'Australia to help staff Ebola treatment centre in Sierra Leone',
-            'description': 'md5:ba36fa5e27e5c9251fd929d339aea4af',
+            'description': 'md5:809ad29c67a05f54eb41f2a105693a67',
        },
    }
--- a/youtube_dl/extractor/academicearth.py
+++ b/youtube_dl/extractor/academicearth.py
@ -1,4 +1,5 @@
 from __future__ import unicode_literals
 import re
 from .common import InfoExtractor
@ -18,15 +19,14 @@ class AcademicEarthCourseIE(InfoExtractor):
    }
    def _real_extract(self, url):
-        m = re.match(self._VALID_URL, url)
+        playlist_id = self._match_id(url)
        playlist_id = m.group('id')
        webpage = self._download_webpage(url, playlist_id)
        title = self._html_search_regex(
-            r'<h1 class="playlist-name"[^>]*?>(.*?)</h1>', webpage, u'title')
+            r'<h1 class="playlist-name"[^>]*?>(.*?)</h1>', webpage, 'title')
        description = self._html_search_regex(
            r'<p class="excerpt"[^>]*?>(.*?)</p>',
-            webpage, u'description', fatal=False)
+            webpage, 'description', fatal=False)
        urls = re.findall(
            r'<li class="lecture-preview">\s*?<a target="_blank" href="([^"]+)">',
            webpage)
--- a/youtube_dl/extractor/addanime.py
+++ b/youtube_dl/extractor/addanime.py
@ -3,19 +3,19 @@ from __future__ import unicode_literals
 import re
 from .common import InfoExtractor
-from ..utils import (
+from ..compat import (
    compat_HTTPError,
    compat_str,
    compat_urllib_parse,
    compat_urllib_parse_urlparse,
-
+)
 from ..utils import (
    ExtractorError,
 )
 class AddAnimeIE(InfoExtractor):
-
+    _VALID_URL = r'^http://(?:\w+\.)?add-anime\.net/watch_video\.php\?(?:.*?)v=(?P<id>[\w_]+)(?:.*)'
    _VALID_URL = r'^http://(?:\w+\.)?add-anime\.net/watch_video\.php\?(?:.*?)v=(?P<video_id>[\w_]+)(?:.*)'
    _TEST = {
        'url': 'http://www.add-anime.net/watch_video.php?v=24MR3YO5SAS9',
        'md5': '72954ea10bc979ab5e2eb288b21425a0',
@ -28,9 +28,9 @@ class AddAnimeIE(InfoExtractor):
    }
    def _real_extract(self, url):
        video_id = self._match_id(url)
        try:
            mobj = re.match(self._VALID_URL, url)
            video_id = mobj.group('video_id')
            webpage = self._download_webpage(url, video_id)
        except ExtractorError as ee:
            if not isinstance(ee.cause, compat_HTTPError) or \
@ -48,7 +48,7 @@ class AddAnimeIE(InfoExtractor):
                r'a\.value = ([0-9]+)[+]([0-9]+)[*]([0-9]+);',
                redir_webpage)
            if av is None:
-                raise ExtractorError(u'Cannot find redirect math task')
+                raise ExtractorError('Cannot find redirect math task')
            av_res = int(av.group(1)) + int(av.group(2)) * int(av.group(3))
            parsed_url = compat_urllib_parse_urlparse(url)
--- a/youtube_dl/extractor/adultswim.py
+++ b/youtube_dl/extractor/adultswim.py
@ -5,6 +5,7 @@ import re
 from .common import InfoExtractor
 class AdultSwimIE(InfoExtractor):
    _VALID_URL = r'https?://video\.adultswim\.com/(?P<path>.+?)(?:\.html)?(?:\?.*)?(?:#.*)?$'
    _TEST = {
--- a/youtube_dl/extractor/allocine.py
+++ b/youtube_dl/extractor/allocine.py
@ -22,7 +22,7 @@ class AllocineIE(InfoExtractor):
            'id': '19546517',
            'ext': 'mp4',
            'title': 'Astérix - Le Domaine des Dieux Teaser VF',
-            'description': 'md5:4a754271d9c6f16c72629a8a993ee884',
+            'description': 'md5:abcd09ce503c6560512c14ebfdb720d2',
            'thumbnail': 're:http://.*\.jpg',
        },
    }, {
--- a/youtube_dl/extractor/aparat.py
+++ b/youtube_dl/extractor/aparat.py
@ -1,5 +1,4 @@
 # coding: utf-8
 from __future__ import unicode_literals
 import re
@ -26,8 +25,7 @@ class AparatIE(InfoExtractor):
    }
    def _real_extract(self, url):
-        m = re.match(self._VALID_URL, url)
+        video_id = self._match_id(url)
        video_id = m.group('id')
        # Note: There is an easier-to-parse configuration at
        # http://www.aparat.com/video/video/config/videohash/%video_id
@ -40,15 +38,15 @@ class AparatIE(InfoExtractor):
        for i, video_url in enumerate(video_urls):
            req = HEADRequest(video_url)
            res = self._request_webpage(
-                req, video_id, note=u'Testing video URL %d' % i, errnote=False)
+                req, video_id, note='Testing video URL %d' % i, errnote=False)
            if res:
                break
        else:
-            raise ExtractorError(u'No working video URLs found')
+            raise ExtractorError('No working video URLs found')
-        title = self._search_regex(r'\s+title:\s*"([^"]+)"', webpage, u'title')
+        title = self._search_regex(r'\s+title:\s*"([^"]+)"', webpage, 'title')
        thumbnail = self._search_regex(
-            r'\s+image:\s*"([^"]+)"', webpage, u'thumbnail', fatal=False)
+            r'\s+image:\s*"([^"]+)"', webpage, 'thumbnail', fatal=False)
        return {
            'id': video_id,
--- a/youtube_dl/extractor/appletrailers.py
+++ b/youtube_dl/extractor/appletrailers.py
@ -70,15 +70,17 @@ class AppleTrailersIE(InfoExtractor):
        uploader_id = mobj.group('company')
        playlist_url = compat_urlparse.urljoin(url, 'includes/playlists/itunes.inc')
        def fix_html(s):
            s = re.sub(r'(?s)<script[^<]*?>.*?</script>', '', s)
            s = re.sub(r'<img ([^<]*?)>', r'<img \1/>', s)
            # The ' in the onClick attributes are not escaped, it couldn't be parsed
            # like: http://trailers.apple.com/trailers/wb/gravity/
            def _clean_json(m):
                return 'iTunes.playURL(%s);' % m.group(1).replace('\'', '&#39;')
            s = re.sub(self._JSON_RE, _clean_json, s)
-            s = '<html>' + s + u'</html>'
+            s = '<html>%s</html>' % s
            return s
        doc = self._download_xml(playlist_url, movie, transform_source=fix_html)
--- a/youtube_dl/extractor/ard.py
+++ b/youtube_dl/extractor/ard.py
@ -192,4 +192,3 @@ class ARDIE(InfoExtractor):
            'upload_date': upload_date,
            'thumbnail': thumbnail,
        }
--- a/youtube_dl/extractor/arte.py
+++ b/youtube_dl/extractor/arte.py
@ -5,13 +5,12 @@ import re
 from .common import InfoExtractor
 from ..utils import (
    ExtractorError,
    find_xpath_attr,
    unified_strdate,
    determine_ext,
    get_element_by_id,
    get_element_by_attribute,
    int_or_none,
    qualities,
 )
 # There are different sources of video in arte.tv, the extraction process
@ -102,79 +101,54 @@ class ArteTVPlus7IE(InfoExtractor):
            'upload_date': unified_strdate(upload_date_str),
            'thumbnail': player_info.get('programImage') or player_info.get('VTU', {}).get('IUR'),
        }
        qfunc = qualities(['HQ', 'MQ', 'EQ', 'SQ'])
-        all_formats = []
+        formats = []
        for format_id, format_dict in player_info['VSR'].items():
-            fmt = dict(format_dict)
+            f = dict(format_dict)
            fmt['format_id'] = format_id
            all_formats.append(fmt)
        # Some formats use the m3u8 protocol
        all_formats = list(filter(lambda f: f.get('videoFormat') != 'M3U8', all_formats))
        def _match_lang(f):
            if f.get('versionCode') is None:
                return True
            # Return true if that format is in the language of the url
            if lang == 'fr':
                l = 'F'
            elif lang == 'de':
                l = 'A'
            else:
                l = lang
            regexes = [r'VO?%s' % l, r'VO?.-ST%s' % l]
            return any(re.match(r, f['versionCode']) for r in regexes)
        # Some formats may not be in the same language as the url
        # TODO: Might want not to drop videos that does not match requested language
        # but to process those formats with lower precedence
        formats = filter(_match_lang, all_formats)
        formats = list(formats)  # in python3 filter returns an iterator
        if not formats:
            # Some videos are only available in the 'Originalversion'
            # they aren't tagged as being in French or German
            # Sometimes there are neither videos of requested lang code
            # nor original version videos available
            # For such cases we just take all_formats as is
            formats = all_formats
            if not formats:
                raise ExtractorError('The formats list is empty')
        if re.match(r'[A-Z]Q', formats[0]['quality']) is not None:
            def sort_key(f):
                return ['HQ', 'MQ', 'EQ', 'SQ'].index(f['quality'])
        else:
            def sort_key(f):
            versionCode = f.get('versionCode')
                if versionCode is None:
                    versionCode = ''
                return (
                    # Sort first by quality
                    int(f.get('height', -1)),
                    int(f.get('bitrate', -1)),
                    # The original version with subtitles has lower relevance
                    re.match(r'VO-ST(F|A)', versionCode) is None,
                    # The version with sourds/mal subtitles has also lower relevance
                    re.match(r'VO?(F|A)-STM\1', versionCode) is None,
                    # Prefer http downloads over m3u8
                    0 if f['url'].endswith('m3u8') else 1,
                )
        formats = sorted(formats, key=sort_key)
        def _format(format_info):
            info = {
                'format_id': format_info['format_id'],
                'format_note': '%s, %s' % (format_info.get('versionCode'), format_info.get('versionLibelle')),
                'width': int_or_none(format_info.get('width')),
                'height': int_or_none(format_info.get('height')),
                'tbr': int_or_none(format_info.get('bitrate')),
            }
            if format_info['mediaType'] == 'rtmp':
                info['url'] = format_info['streamer']
                info['play_path'] = 'mp4:' + format_info['url']
                info['ext'] = 'flv'
            else:
                info['url'] = format_info['url']
                info['ext'] = determine_ext(info['url'])
            return info
        info_dict['formats'] = [_format(f) for f in formats]
            langcode = {
                'fr': 'F',
                'de': 'A',
            }.get(lang, lang)
            lang_rexs = [r'VO?%s' % langcode, r'VO?.-ST%s' % langcode]
            lang_pref = (
                None if versionCode is None else (
                    10 if any(re.match(r, versionCode) for r in lang_rexs)
                    else -10))
            source_pref = 0
            if versionCode is not None:
                # The original version with subtitles has lower relevance
                if re.match(r'VO-ST(F|A)', versionCode):
                    source_pref -= 10
                # The version with sourds/mal subtitles has also lower relevance
                elif re.match(r'VO?(F|A)-STM\1', versionCode):
                    source_pref -= 9
            format = {
                'format_id': format_id,
                'preference': -10 if f.get('videoFormat') == 'M3U8' else None,
                'language_preference': lang_pref,
                'format_note': '%s, %s' % (f.get('versionCode'), f.get('versionLibelle')),
                'width': int_or_none(f.get('width')),
                'height': int_or_none(f.get('height')),
                'tbr': int_or_none(f.get('bitrate')),
                'quality': qfunc(f['quality']),
                'source_preference': source_pref,
            }
            if f.get('mediaType') == 'rtmp':
                format['url'] = f['streamer']
                format['play_path'] = 'mp4:' + f['url']
                format['ext'] = 'flv'
            else:
                format['url'] = f['url']
            formats.append(format)
        self._sort_formats(formats)
        info_dict['formats'] = formats
        return info_dict
--- a/youtube_dl/extractor/bambuser.py
+++ b/youtube_dl/extractor/bambuser.py
@ -18,7 +18,7 @@ class BambuserIE(InfoExtractor):
    _TEST = {
        'url': 'http://bambuser.com/v/4050584',
        # MD5 seems to be flaky, see https://travis-ci.org/rg3/youtube-dl/jobs/14051016#L388
-        #u'md5': 'fba8f7693e48fd4e8641b3fd5539a641',
+        # 'md5': 'fba8f7693e48fd4e8641b3fd5539a641',
        'info_dict': {
            'id': '4050584',
            'ext': 'flv',
@ -73,7 +73,8 @@ class BambuserChannelIE(InfoExtractor):
        urls = []
        last_id = ''
        for i in itertools.count(1):
-            req_url = ('http://bambuser.com/xhr-api/index.php?username={user}'
+            req_url = (
                'http://bambuser.com/xhr-api/index.php?username={user}'
                '&sort=created&access_mode=0%2C1%2C2&limit={count}'
                '&method=broadcast&format=json&vid_older_than={last}'
            ).format(user=user, count=self._STEP, last=last_id)
--- a/youtube_dl/extractor/bandcamp.py
+++ b/youtube_dl/extractor/bandcamp.py
@ -110,20 +110,25 @@ class BandcampAlbumIE(InfoExtractor):
        'url': 'http://blazo.bandcamp.com/album/jazz-format-mixtape-vol-1',
        'playlist': [
            {
                'file': '1353101989.mp3',
                'md5': '39bc1eded3476e927c724321ddf116cf',
                'info_dict': {
                    'id': '1353101989',
                    'ext': 'mp3',
                    'title': 'Intro',
                }
            },
            {
                'file': '38097443.mp3',
                'md5': '1a2c32e2691474643e912cc6cd4bffaa',
                'info_dict': {
                    'id': '38097443',
                    'ext': 'mp3',
                    'title': 'Kero One - Keep It Alive (Blazo remix)',
                }
            },
        ],
        'info_dict': {
            'title': 'Jazz Format Mixtape vol.1',
        },
        'params': {
            'playlistend': 2
        },
--- a/youtube_dl/extractor/bliptv.py
+++ b/youtube_dl/extractor/bliptv.py
@ -71,11 +71,12 @@ class BlipTVIE(SubtitlesInfoExtractor):
        mobj = re.match(self._VALID_URL, url)
        lookup_id = mobj.group('lookup_id')
-        # See https://github.com/rg3/youtube-dl/issues/857
+        # See https://github.com/rg3/youtube-dl/issues/857 and
        # https://github.com/rg3/youtube-dl/issues/4197
        if lookup_id:
            info_page = self._download_webpage(
                'http://blip.tv/play/%s.x?p=1' % lookup_id, lookup_id, 'Resolving lookup id')
-            video_id = self._search_regex(r'data-episode-id="([0-9]+)', info_page, 'video_id')
+            video_id = self._search_regex(r'config\.id\s*=\s*"([0-9]+)', info_page, 'video_id')
        else:
            video_id = mobj.group('id')
@ -165,9 +166,17 @@ class BlipTVIE(SubtitlesInfoExtractor):
 class BlipTVUserIE(InfoExtractor):
-    _VALID_URL = r'(?:(?:(?:https?://)?(?:\w+\.)?blip\.tv/)|bliptvuser:)(?!api\.swf)([^/]+)/*$'
+    _VALID_URL = r'(?:(?:https?://(?:\w+\.)?blip\.tv/)|bliptvuser:)(?!api\.swf)([^/]+)/*$'
    _PAGE_SIZE = 12
    IE_NAME = 'blip.tv:user'
    _TEST = {
        'url': 'http://blip.tv/actone',
        'info_dict': {
            'id': 'actone',
            'title': 'Act One: The Series',
        },
        'playlist_count': 5,
    }
    def _real_extract(self, url):
        mobj = re.match(self._VALID_URL, url)
@ -178,6 +187,7 @@ class BlipTVUserIE(InfoExtractor):
        page = self._download_webpage(url, username, 'Downloading user page')
        mobj = re.search(r'data-users-id="([^"]+)"', page)
        page_base = page_base % mobj.group(1)
        title = self._og_search_title(page)
        # Download video ids using BlipTV Ajax calls. Result size per
        # query is limited (currently to 12 videos) so we need to query
@ -214,4 +224,5 @@ class BlipTVUserIE(InfoExtractor):
        urls = ['http://blip.tv/%s' % video_id for video_id in video_ids]
        url_entries = [self.url_result(vurl, 'BlipTV') for vurl in urls]
-        return [self.playlist_result(url_entries, playlist_title=username)]
+        return self.playlist_result(
            url_entries, playlist_title=title, playlist_id=username)
--- a/youtube_dl/extractor/bpb.py
+++ b/youtube_dl/extractor/bpb.py
@ -0,0 +1,37 @@
 # coding: utf-8
 from __future__ import unicode_literals
 from .common import InfoExtractor
 class BpbIE(InfoExtractor):
    IE_DESC = 'Bundeszentrale für politische Bildung'
    _VALID_URL = r'http://www\.bpb\.de/mediathek/(?P<id>[0-9]+)/'
    _TEST = {
        'url': 'http://www.bpb.de/mediathek/297/joachim-gauck-zu-1989-und-die-erinnerung-an-die-ddr',
        'md5': '0792086e8e2bfbac9cdf27835d5f2093',
        'info_dict': {
            'id': '297',
            'ext': 'mp4',
            'title': 'Joachim Gauck zu 1989 und die Erinnerung an die DDR',
            'description': 'Joachim Gauck, erster Beauftragter für die Stasi-Unterlagen, spricht auf dem Geschichtsforum über die friedliche Revolution 1989 und eine "gewisse Traurigkeit" im Umgang mit der DDR-Vergangenheit.'
        }
    }
    def _real_extract(self, url):
        video_id = self._match_id(url)
        webpage = self._download_webpage(url, video_id)
        title = self._html_search_regex(
            r'<h2 class="white">(.*?)</h2>', webpage, 'title')
        video_url = self._html_search_regex(
            r'(http://film\.bpb\.de/player/dokument_[0-9]+\.mp4)',
            webpage, 'video URL')
        return {
            'id': video_id,
            'url': video_url,
            'title': title,
            'description': self._og_search_description(webpage),
        }
--- a/youtube_dl/extractor/brightcove.py
+++ b/youtube_dl/extractor/brightcove.py
@ -14,6 +14,7 @@ from ..utils import (
    compat_str,
    compat_urllib_request,
    compat_parse_qs,
    compat_urllib_parse_urlparse,
    determine_ext,
    ExtractorError,
@ -23,7 +24,7 @@ from ..utils import (
 class BrightcoveIE(InfoExtractor):
-    _VALID_URL = r'https?://.*brightcove\.com/(services|viewer).*\?(?P<query>.*)'
+    _VALID_URL = r'https?://.*brightcove\.com/(services|viewer).*?\?(?P<query>.*)'
    _FEDERATED_URL_TEMPLATE = 'http://c.brightcove.com/services/viewer/htmlFederated?%s'
    _TESTS = [
@ -110,6 +111,8 @@ class BrightcoveIE(InfoExtractor):
                            lambda m: m.group(1) + '/>', object_str)
        # Fix up some stupid XML, see https://github.com/rg3/youtube-dl/issues/1608
        object_str = object_str.replace('<--', '<!--')
        # remove namespace to simplify extraction
        object_str = re.sub(r'(<object[^>]*)(xmlns=".*?")', r'\1', object_str)
        object_str = fix_xml_ampersands(object_str)
        object_doc = xml.etree.ElementTree.fromstring(object_str.encode('utf-8'))
@ -218,7 +221,7 @@ class BrightcoveIE(InfoExtractor):
        webpage = self._download_webpage(req, video_id)
        error_msg = self._html_search_regex(
-            r"<h1>We're sorry.</h1>\s*<p>(.*?)</p>", webpage,
+            r"<h1>We're sorry.</h1>([\s\n]*<p>.*?</p>)+", webpage,
            'error message', default=None)
        if error_msg is not None:
            raise ExtractorError(
@ -260,9 +263,17 @@ class BrightcoveIE(InfoExtractor):
            formats = []
            for rend in renditions:
                url = rend['defaultURL']
                if not url:
                    continue
                if rend['remote']:
-                    # This type of renditions are served through akamaihd.net,
+                    url_comp = compat_urllib_parse_urlparse(url)
-                    # but they don't use f4m manifests
+                    if url_comp.path.endswith('.m3u8'):
                        formats.extend(
                            self._extract_m3u8_formats(url, info['id'], 'mp4'))
                        continue
                    elif 'akamaihd.net' in url_comp.netloc:
                        # This type of renditions are served through
                        # akamaihd.net, but they don't use f4m manifests
                        url = url.replace('control/', '') + '?&v=3.3.0&fp=13&r=FEEFJ&g=RTSJIMBMPFPB'
                        ext = 'flv'
                else:
--- a/youtube_dl/extractor/byutv.py
+++ b/youtube_dl/extractor/byutv.py
@ -10,12 +10,12 @@ from ..utils import ExtractorError
 class BYUtvIE(InfoExtractor):
    _VALID_URL = r'^https?://(?:www\.)?byutv.org/watch/[0-9a-f-]+/(?P<video_id>[^/?#]+)'
    _TEST = {
-        'url': 'http://www.byutv.org/watch/44e80f7b-e3ba-43ba-8c51-b1fd96c94a79/granite-flats-talking',
+        'url': 'http://www.byutv.org/watch/6587b9a3-89d2-42a6-a7f7-fd2f81840a7d/studio-c-season-5-episode-5',
        'info_dict': {
-            'id': 'granite-flats-talking',
+            'id': 'studio-c-season-5-episode-5',
            'ext': 'mp4',
-            'description': 'md5:4e9a7ce60f209a33eca0ac65b4918e1c',
+            'description': 'md5:5438d33774b6bdc662f9485a340401cc',
-            'title': 'Talking',
+            'title': 'Season 5 Episode 5',
            'thumbnail': 're:^https?://.*promo.*'
        },
        'params': {
--- a/youtube_dl/extractor/cbs.py
+++ b/youtube_dl/extractor/cbs.py
@ -45,4 +45,4 @@ class CBSIE(InfoExtractor):
        real_id = self._search_regex(
            r"video\.settings\.pid\s*=\s*'([^']+)';",
            webpage, 'real video ID')
-        return self.url_result(u'theplatform:%s' % real_id)
+        return self.url_result('theplatform:%s' % real_id)
--- a/youtube_dl/extractor/channel9.py
+++ b/youtube_dl/extractor/channel9.py
@ -5,6 +5,7 @@ import re
 from .common import InfoExtractor
 from ..utils import ExtractorError
 class Channel9IE(InfoExtractor):
    '''
    Common extractor for channel9.msdn.com.
@ -27,7 +28,7 @@ class Channel9IE(InfoExtractor):
                'title': 'Developer Kick-Off Session: Stuff We Love',
                'description': 'md5:c08d72240b7c87fcecafe2692f80e35f',
                'duration': 4576,
-                'thumbnail': 'http://media.ch9.ms/ch9/9d51/03902f2d-fc97-4d3c-b195-0bfe15a19d51/KOS002_220.jpg',
+                'thumbnail': 'http://video.ch9.ms/ch9/9d51/03902f2d-fc97-4d3c-b195-0bfe15a19d51/KOS002_220.jpg',
                'session_code': 'KOS002',
                'session_day': 'Day 1',
                'session_room': 'Arena 1A',
@ -43,7 +44,7 @@ class Channel9IE(InfoExtractor):
                'title': 'Self-service BI with Power BI - nuclear testing',
                'description': 'md5:d1e6ecaafa7fb52a2cacdf9599829f5b',
                'duration': 1540,
-                'thumbnail': 'http://media.ch9.ms/ch9/87e1/0300391f-a455-4c72-bec3-4422f19287e1/selfservicenuk_512.jpg',
+                'thumbnail': 'http://video.ch9.ms/ch9/87e1/0300391f-a455-4c72-bec3-4422f19287e1/selfservicenuk_512.jpg',
                'authors': ['Mike Wilmot'],
            },
        }
@ -115,7 +116,7 @@ class Channel9IE(InfoExtractor):
        return self._html_search_meta('description', html, 'description')
    def _extract_duration(self, html):
-        m = re.search(r'data-video_duration="(?P<hours>\d{2}):(?P<minutes>\d{2}):(?P<seconds>\d{2})"', html)
+        m = re.search(r'"length": *"(?P<hours>\d{2}):(?P<minutes>\d{2}):(?P<seconds>\d{2})"', html)
        return ((int(m.group('hours')) * 60 * 60) + (int(m.group('minutes')) * 60) + int(m.group('seconds'))) if m else None
    def _extract_slides(self, html):
@ -187,7 +188,8 @@ class Channel9IE(InfoExtractor):
        view_count = self._extract_view_count(html)
        comment_count = self._extract_comment_count(html)
-        common = {'_type': 'video',
+        common = {
            '_type': 'video',
            'id': content_path,
            'description': description,
            'thumbnail': thumbnail,
@ -258,16 +260,17 @@ class Channel9IE(InfoExtractor):
        webpage = self._download_webpage(url, content_path, 'Downloading web page')
-        page_type_m = re.search(r'<meta name="Search.PageType" content="(?P<pagetype>[^"]+)"/>', webpage)
+        page_type_m = re.search(r'<meta name="WT.entryid" content="(?P<pagetype>[^:]+)[^"]+"/>', webpage)
-        if page_type_m is None:
+        if page_type_m is not None:
            raise ExtractorError('Search.PageType not found, don\'t know how to process this page', expected=True)
            page_type = page_type_m.group('pagetype')
-        if page_type == 'List':         # List page, may contain list of 'item'-like objects
+            if page_type == 'Entry':      # Any 'item'-like page, may contain downloadable content
            return self._extract_list(content_path)
        elif page_type == 'Entry.Item': # Any 'item'-like page, may contain downloadable content
                return self._extract_entry_item(webpage, content_path)
            elif page_type == 'Session':  # Event session page, may contain downloadable content
                return self._extract_session(webpage, content_path)
            elif page_type == 'Event':
                return self._extract_list(content_path)
            else:
-            raise ExtractorError('Unexpected Search.PageType %s' % page_type, expected=True)
+                raise ExtractorError('Unexpected WT.entryid %s' % page_type, expected=True)
        else:  # Assuming list
            return self._extract_list(content_path)
--- a/youtube_dl/extractor/cinemassacre.py
+++ b/youtube_dl/extractor/cinemassacre.py
@ -42,11 +42,12 @@ class CinemassacreIE(InfoExtractor):
        webpage = self._download_webpage(url, display_id)
        video_date = mobj.group('date_Y') + mobj.group('date_m') + mobj.group('date_d')
-        mobj = re.search(r'src="(?P<embed_url>http://player\.screenwavemedia\.com/play/[a-zA-Z]+\.php\?[^"]*\bid=(?:Cinemassacre-)?(?P<video_id>.+?))"', webpage)
+        mobj = re.search(r'src="(?P<embed_url>http://player\.screenwavemedia\.com/play/[a-zA-Z]+\.php\?[^"]*\bid=(?P<full_video_id>(?:Cinemassacre-)?(?P<video_id>.+?)))"', webpage)
        if not mobj:
            raise ExtractorError('Can\'t extract embed url and video id')
        playerdata_url = mobj.group('embed_url')
        video_id = mobj.group('video_id')
        full_video_id = mobj.group('full_video_id')
        video_title = self._html_search_regex(
            r'<title>(?P<title>.+?)\|', webpage, 'title')
@ -60,10 +61,21 @@ class CinemassacreIE(InfoExtractor):
        vidurl = self._search_regex(
            r'\'vidurl\'\s*:\s*"([^\']+)"', playerdata, 'vidurl').replace('\\/', '/')
-        videolist_url = self._search_regex(
+        videolist_url = None
            r"file\s*:\s*'(http.+?/jwplayer\.smil)'", playerdata, 'jwplayer.smil')
        videolist = self._download_xml(videolist_url, video_id, 'Downloading videolist XML')
        mobj = re.search(r"'videoserver'\s*:\s*'(?P<videoserver>[^']+)'", playerdata)
        if mobj:
            videoserver = mobj.group('videoserver')
            mobj = re.search(r'\'vidid\'\s*:\s*"(?P<vidid>[^\']+)"', playerdata)
            vidid = mobj.group('vidid') if mobj else full_video_id
            videolist_url = 'http://%s/vod/smil:%s.smil/jwplayer.smil' % (videoserver, vidid)
        else:
            mobj = re.search(r"file\s*:\s*'(?P<smil>http.+?/jwplayer\.smil)'", playerdata)
            if mobj:
                videolist_url = mobj.group('smil')
        if videolist_url:
            videolist = self._download_xml(videolist_url, video_id, 'Downloading videolist XML')
            formats = []
            baseurl = vidurl[:vidurl.rfind('/') + 1]
            for video in videolist.findall('.//video'):
@ -91,6 +103,10 @@ class CinemassacreIE(InfoExtractor):
                    })
                formats.append(format)
            self._sort_formats(formats)
        else:
            formats = [{
                'url': vidurl,
            }]
        return {
            'id': video_id,
--- a/youtube_dl/extractor/clipfish.py
+++ b/youtube_dl/extractor/clipfish.py
@ -24,7 +24,7 @@ class ClipfishIE(InfoExtractor):
            'title': 'FIFA 14 - E3 2013 Trailer',
            'duration': 82,
        },
-        u'skip': 'Blocked in the US'
+        'skip': 'Blocked in the US'
    }
    def _real_extract(self, url):
@ -34,7 +34,7 @@ class ClipfishIE(InfoExtractor):
        info_url = ('http://www.clipfish.de/devxml/videoinfo/%s?ts=%d' %
                    (video_id, int(time.time())))
        doc = self._download_xml(
-            info_url, video_id, note=u'Downloading info page')
+            info_url, video_id, note='Downloading info page')
        title = doc.find('title').text
        video_url = doc.find('filename').text
        if video_url is None:
--- a/youtube_dl/extractor/clipsyndicate.py
+++ b/youtube_dl/extractor/clipsyndicate.py
@ -39,6 +39,7 @@ class ClipsyndicateIE(InfoExtractor):
            transform_source=fix_xml_ampersands)
        track_doc = pdoc.find('trackList/track')
        def find_param(name):
            node = find_xpath_attr(track_doc, './/param', 'name', name)
            if node is not None:
--- a/youtube_dl/extractor/cloudy.py
+++ b/youtube_dl/extractor/cloudy.py
@ -4,14 +4,16 @@ from __future__ import unicode_literals
 import re
 from .common import InfoExtractor
-from ..utils import (
+from ..compat import (
    ExtractorError,
    compat_parse_qs,
    compat_urllib_parse,
    remove_end,
    HEADRequest,
    compat_HTTPError,
 )
 from ..utils import (
    ExtractorError,
    HEADRequest,
    remove_end,
 )
 class CloudyIE(InfoExtractor):
--- a/youtube_dl/extractor/cnn.py
+++ b/youtube_dl/extractor/cnn.py
@ -16,20 +16,21 @@ class CNNIE(InfoExtractor):
    _TESTS = [{
        'url': 'http://edition.cnn.com/video/?/video/sports/2013/06/09/nadal-1-on-1.cnn',
        'file': 'sports_2013_06_09_nadal-1-on-1.cnn.mp4',
        'md5': '3e6121ea48df7e2259fe73a0628605c4',
        'info_dict': {
            'id': 'sports_2013_06_09_nadal-1-on-1.cnn',
            'ext': 'mp4',
            'title': 'Nadal wins 8th French Open title',
            'description': 'World Sport\'s Amanda Davies chats with 2013 French Open champion Rafael Nadal.',
            'duration': 135,
            'upload_date': '20130609',
        },
-    },
+    }, {
    {
        "url": "http://edition.cnn.com/video/?/video/us/2013/08/21/sot-student-gives-epic-speech.georgia-institute-of-technology&utm_source=feedburner&utm_medium=feed&utm_campaign=Feed%3A+rss%2Fcnn_topstories+%28RSS%3A+Top+Stories%29",
        "file": "us_2013_08_21_sot-student-gives-epic-speech.georgia-institute-of-technology.mp4",
        "md5": "b5cc60c60a3477d185af8f19a2a26f4e",
        "info_dict": {
            'id': 'us/2013/08/21/sot-student-gives-epic-speech.georgia-institute-of-technology',
            'ext': 'mp4',
            "title": "Student's epic speech stuns new freshmen",
            "description": "A Georgia Tech student welcomes the incoming freshmen with an epic speech backed by music from \"2001: A Space Odyssey.\"",
            "upload_date": "20130821",
--- a/youtube_dl/extractor/collegehumor.py
+++ b/youtube_dl/extractor/collegehumor.py
@ -10,7 +10,8 @@ from ..utils import int_or_none
 class CollegeHumorIE(InfoExtractor):
    _VALID_URL = r'^(?:https?://)?(?:www\.)?collegehumor\.com/(video|embed|e)/(?P<videoid>[0-9]+)/?(?P<shorttitle>.*)$'
-    _TESTS = [{
+    _TESTS = [
        {
            'url': 'http://www.collegehumor.com/video/6902724/comic-con-cosplay-catastrophe',
            'md5': 'dcc0f5c1c8be98dc33889a191f4c26bd',
            'info_dict': {
@ -21,8 +22,7 @@ class CollegeHumorIE(InfoExtractor):
                'age_limit': 13,
                'duration': 187,
            },
-    },
+        }, {
    {
            'url': 'http://www.collegehumor.com/video/3505939/font-conference',
            'md5': '72fa701d8ef38664a4dbb9e2ab721816',
            'info_dict': {
@ -33,9 +33,8 @@ class CollegeHumorIE(InfoExtractor):
                'age_limit': 10,
                'duration': 179,
            },
-    },
+        }, {
            # embedded youtube video
    {
            'url': 'http://www.collegehumor.com/embed/6950306',
            'info_dict': {
                'id': 'Z-bao9fg6Yc',
--- a/youtube_dl/extractor/comedycentral.py
+++ b/youtube_dl/extractor/comedycentral.py
@ -2,7 +2,6 @@ from __future__ import unicode_literals
 import re
 from .common import InfoExtractor
 from .mtv import MTVServicesInfoExtractor
 from ..utils import (
    compat_str,
@ -31,7 +30,7 @@ class ComedyCentralIE(MTVServicesInfoExtractor):
    }
-class ComedyCentralShowsIE(InfoExtractor):
+class ComedyCentralShowsIE(MTVServicesInfoExtractor):
    IE_DESC = 'The Daily Show / The Colbert Report'
    # urls can be abbreviations like :thedailyshow or :colbert
    # urls for episodes like:
@ -109,18 +108,8 @@ class ComedyCentralShowsIE(InfoExtractor):
        '400': (384, 216),
    }
    @staticmethod
    def _transform_rtmp_url(rtmp_video_url):
        m = re.match(r'^rtmpe?://.*?/(?P<finalid>gsp\.comedystor/.*)$', rtmp_video_url)
        if not m:
            raise ExtractorError('Cannot transform RTMP url')
        base = 'http://mtvnmobile.vo.llnwd.net/kip0/_pxn=1+_pxI0=Ripod-h264+_pxL0=undefined+_pxM0=+_pxK=18639+_pxE=mp4/44620/mtvnorigin/'
        return base + m.group('finalid')
    def _real_extract(self, url):
-        mobj = re.match(self._VALID_URL, url, re.VERBOSE)
+        mobj = re.match(self._VALID_URL, url)
        if mobj is None:
            raise ExtractorError('Invalid URL: %s' % url)
        if mobj.group('shortname'):
            if mobj.group('shortname') in ('tds', 'thedailyshow'):
@ -212,9 +201,6 @@ class ComedyCentralShowsIE(InfoExtractor):
                    'ext': self._video_extensions.get(format, 'mp4'),
                    'height': h,
                    'width': w,
                    'format_note': 'HTTP 400 at the moment (patches welcome!)',
                    'preference': -100,
                })
                formats.append({
                    'format_id': 'rtmp-%s' % format,
--- a/youtube_dl/extractor/common.py
+++ b/youtube_dl/extractor/common.py
@ -12,13 +12,14 @@ import sys
 import time
 import xml.etree.ElementTree
-from ..utils import (
+from ..compat import (
    compat_http_client,
    compat_urllib_error,
    compat_urllib_parse_urlparse,
    compat_urlparse,
    compat_str,
-
+)
 from ..utils import (
    clean_html,
    compiled_regex_type,
    ExtractorError,
@ -42,7 +43,11 @@ class InfoExtractor(object):
    information possibly downloading the video to the file system, among
    other possible outcomes.
-    The dictionaries must include the following fields:
+    The type field determines the the type of the result.
    By far the most common value (and the default if _type is missing) is
    "video", which indicates a single video.
    For a video, the dictionaries must include the following fields:
    id:             Video identifier.
    title:          Video title, unescaped.
@ -86,6 +91,11 @@ class InfoExtractor(object):
                                 by this field, regardless of all other values.
                                 -1 for default (order by other properties),
                                 -2 or smaller for less than default.
                    * language_preference  Is this in the correct requested
                                 language?
                                 10 if it's what the URL is about,
                                 -1 for default (don't know),
                                 -10 otherwise, other values reserved for now.
                    * quality    Order number of the video quality of this
                                 format, irrespective of the file format.
                                 -1 for default (order by other properties),
@ -145,6 +155,38 @@ class InfoExtractor(object):
    Unless mentioned otherwise, None is equivalent to absence of information.
    _type "playlist" indicates multiple videos.
    There must be a key "entries", which is a list or a PagedList object, each
    element of which is a valid dictionary under this specfication.
    Additionally, playlists can have "title" and "id" attributes with the same
    semantics as videos (see above).
    _type "multi_video" indicates that there are multiple videos that
    form a single show, for examples multiple acts of an opera or TV episode.
    It must have an entries key like a playlist and contain all the keys
    required for a video at the same time.
    _type "url" indicates that the video must be extracted from another
    location, possibly by a different extractor. Its only required key is:
    "url" - the next URL to extract.
    Additionally, it may have properties believed to be identical to the
    resolved entity, for example "title" if the title of the referred video is
    known ahead of time.
    _type "url_transparent" entities have the same specification as "url", but
    indicate that the given additional information is more precise than the one
    associated with the resolved URL.
    This is useful when a site employs a video service that hosts the video and
    its technical metadata, but that video service does not embed a useful
    title, description etc.
    Subclasses of this one should re-define the _real_initialize() and
    _real_extract() methods and define a _VALID_URL regexp.
    Probably, they should also be added to the list of extractors.
@ -254,9 +296,11 @@ class InfoExtractor(object):
        content = self._webpage_read_content(urlh, url_or_request, video_id, note, errnote, fatal)
        return (content, urlh)
-    def _webpage_read_content(self, urlh, url_or_request, video_id, note=None, errnote=None, fatal=True):
+    def _webpage_read_content(self, urlh, url_or_request, video_id, note=None, errnote=None, fatal=True, prefix=None):
        content_type = urlh.headers.get('Content-Type', '')
        webpage_bytes = urlh.read()
        if prefix is not None:
            webpage_bytes = prefix + webpage_bytes
        m = re.match(r'[a-zA-Z0-9_.-]+/[a-zA-Z0-9_.-]+\s*;\s*charset=(.+)', content_type)
        if m:
            encoding = m.group(1)
@ -392,6 +436,7 @@ class InfoExtractor(object):
        if video_id is not None:
            video_info['id'] = video_id
        return video_info
    @staticmethod
    def playlist_result(entries, playlist_id=None, playlist_title=None):
        """Returns a playlist"""
@ -403,7 +448,7 @@ class InfoExtractor(object):
            video_info['title'] = playlist_title
        return video_info
-    def _search_regex(self, pattern, string, name, default=_NO_DEFAULT, fatal=True, flags=0):
+    def _search_regex(self, pattern, string, name, default=_NO_DEFAULT, fatal=True, flags=0, group=None):
        """
        Perform a regex search on the given string, using a single or a list of
        patterns returning the first matching group.
@ -424,8 +469,11 @@ class InfoExtractor(object):
            _name = name
        if mobj:
            if group is None:
                # return the first matching group
                return next(g for g in mobj.groups() if g is not None)
            else:
                return mobj.group(group)
        elif default is not _NO_DEFAULT:
            return default
        elif fatal:
@ -435,11 +483,11 @@ class InfoExtractor(object):
                                            'please report this issue on http://yt-dl.org/bug' % _name)
            return None
-    def _html_search_regex(self, pattern, string, name, default=_NO_DEFAULT, fatal=True, flags=0):
+    def _html_search_regex(self, pattern, string, name, default=_NO_DEFAULT, fatal=True, flags=0, group=None):
        """
        Like _search_regex, but strips HTML tags and unescapes entities.
        """
-        res = self._search_regex(pattern, string, name, default, fatal, flags)
+        res = self._search_regex(pattern, string, name, default, fatal, flags, group)
        if res:
            return clean_html(res).strip()
        else:
@ -533,9 +581,9 @@ class InfoExtractor(object):
            display_name = name
        return self._html_search_regex(
            r'''(?ix)<meta
-                    (?=[^>]+(?:itemprop|name|property)=["\']?%s["\']?)
+                    (?=[^>]+(?:itemprop|name|property)=(["\']?)%s\1)
-                    [^>]+content=["\']([^"\']+)["\']''' % re.escape(name),
+                    [^>]+content=(["\'])(?P<content>.*?)\1''' % re.escape(name),
-            html, display_name, fatal=fatal, **kwargs)
+            html, display_name, fatal=fatal, group='content', **kwargs)
    def _dc_search_uploader(self, html):
        return self._html_search_meta('dc.creator', html, 'uploader')
@ -611,6 +659,7 @@ class InfoExtractor(object):
            return (
                preference,
                f.get('language_preference') if f.get('language_preference') is not None else -1,
                f.get('quality') if f.get('quality') is not None else -1,
                f.get('height') if f.get('height') is not None else -1,
                f.get('width') if f.get('width') is not None else -1,
--- a/youtube_dl/extractor/crunchyroll.py
+++ b/youtube_dl/extractor/crunchyroll.py
@ -17,7 +17,6 @@ from ..utils import (
    bytes_to_intlist,
    intlist_to_bytes,
    unified_strdate,
    clean_html,
    urlencode_postdata,
 )
 from ..aes import (
@ -70,11 +69,9 @@ class CrunchyrollIE(SubtitlesInfoExtractor):
        login_request.add_header('Content-Type', 'application/x-www-form-urlencoded')
        self._download_webpage(login_request, None, False, 'Wrong login info')
    def _real_initialize(self):
        self._login()
    def _decrypt_subtitles(self, data, iv, id):
        data = bytes_to_intlist(data)
        iv = bytes_to_intlist(iv)
@ -100,8 +97,10 @@ class CrunchyrollIE(SubtitlesInfoExtractor):
            return shaHash + [0] * 12
        key = obfuscate_key(id)
        class Counter:
            __value = iv
            def next_value(self):
                temp = self.__value
                self.__value = inc(self.__value)
@ -249,7 +248,8 @@ Format: Layer, Start, End, Style, Name, MarginL, MarginR, MarginV, Effect, Text
        subtitles = {}
        sub_format = self._downloader.params.get('subtitlesformat', 'srt')
        for sub_id, sub_name in re.findall(r'\?ssid=([0-9]+)" title="([^"]+)', webpage):
-            sub_page = self._download_webpage('http://www.crunchyroll.com/xml/?req=RpcApiSubtitle_GetXml&subtitle_script_id='+sub_id,\
+            sub_page = self._download_webpage(
                'http://www.crunchyroll.com/xml/?req=RpcApiSubtitle_GetXml&subtitle_script_id=' + sub_id,
                video_id, note='Downloading subtitles for ' + sub_name)
            id = self._search_regex(r'id=\'([0-9]+)', sub_page, 'subtitle_id', fatal=False)
            iv = self._search_regex(r'<iv>([^<]+)', sub_page, 'subtitle_iv', fatal=False)
@ -265,8 +265,6 @@ Format: Layer, Start, End, Style, Name, MarginL, MarginR, MarginV, Effect, Text
            if not lang_code:
                continue
            sub_root = xml.etree.ElementTree.fromstring(subtitle)
            if not sub_root:
                subtitles[lang_code] = ''
            if sub_format == 'ass':
                subtitles[lang_code] = self._convert_subtitles_to_ass(sub_root)
            else:
--- a/youtube_dl/extractor/dailymotion.py
+++ b/youtube_dl/extractor/dailymotion.py
@ -18,6 +18,7 @@ from ..utils import (
    unescapeHTML,
 )
 class DailymotionBaseInfoExtractor(InfoExtractor):
    @staticmethod
    def _build_request(url):
@ -27,6 +28,7 @@ class DailymotionBaseInfoExtractor(InfoExtractor):
        request.add_header('Cookie', 'ff=off')
        return request
 class DailymotionIE(DailymotionBaseInfoExtractor, SubtitlesInfoExtractor):
    """Information Extractor for Dailymotion"""
@ -94,7 +96,7 @@ class DailymotionIE(DailymotionBaseInfoExtractor, SubtitlesInfoExtractor):
        # It may just embed a vevo video:
        m_vevo = re.search(
-            r'<link rel="video_src" href="[^"]*?vevo.com[^"]*?videoId=(?P<id>[\w]*)',
+            r'<link rel="video_src" href="[^"]*?vevo.com[^"]*?video=(?P<id>[\w]*)',
            webpage)
        if m_vevo is not None:
            vevo_id = m_vevo.group('id')
--- a/youtube_dl/extractor/dropbox.py
+++ b/youtube_dl/extractor/dropbox.py
@ -5,20 +5,21 @@ import os.path
 import re
 from .common import InfoExtractor
-from ..utils import compat_urllib_parse_unquote, url_basename
+from ..compat import compat_urllib_parse_unquote
 from ..utils import url_basename
 class DropboxIE(InfoExtractor):
    _VALID_URL = r'https?://(?:www\.)?dropbox[.]com/sh?/(?P<id>[a-zA-Z0-9]{15})/.*'
-    _TESTS = [{
+    _TESTS = [
        {
            'url': 'https://www.dropbox.com/s/nelirfsxnmcfbfh/youtube-dl%20test%20video%20%27%C3%A4%22BaW_jenozKc.mp4?dl=0',
            'info_dict': {
                'id': 'nelirfsxnmcfbfh',
                'ext': 'mp4',
                'title': 'youtube-dl test video \'ä"BaW_jenozKc'
            }
-    },
+        }, {
    {
            'url': 'https://www.dropbox.com/sh/662glsejgzoj9sr/AAByil3FGH9KFNZ13e08eSa1a/Pregame%20Ceremony%20Program%20PA%2020140518.m4v',
            'only_matching': True,
        },
--- a/youtube_dl/extractor/eighttracks.py
+++ b/youtube_dl/extractor/eighttracks.py
@ -125,7 +125,7 @@ class EightTracksIE(InfoExtractor):
            info = {
                'id': compat_str(track_data['id']),
                'url': track_data['track_file_stream_url'],
-                'title': track_data['performer'] + u' - ' + track_data['name'],
+                'title': track_data['performer'] + ' - ' + track_data['name'],
                'raw_title': track_data['name'],
                'uploader_id': data['user']['login'],
                'ext': 'm4a',
--- a/youtube_dl/extractor/eporner.py
+++ b/youtube_dl/extractor/eporner.py
@ -20,7 +20,7 @@ class EpornerIE(InfoExtractor):
            'display_id': 'Infamous-Tiffany-Teen-Strip-Tease-Video',
            'ext': 'mp4',
            'title': 'Infamous Tiffany Teen Strip Tease Video',
-            'duration': 194,
+            'duration': 1838,
            'view_count': int,
            'age_limit': 18,
        }
@ -57,9 +57,7 @@ class EpornerIE(InfoExtractor):
            formats.append(fmt)
        self._sort_formats(formats)
-        duration = parse_duration(self._search_regex(
+        duration = parse_duration(self._html_search_meta('duration', webpage))
            r'class="mbtim">([0-9:]+)</div>', webpage, 'duration',
            fatal=False))
        view_count = str_to_int(self._search_regex(
            r'id="cinemaviews">\s*([0-9,]+)\s*<small>views',
            webpage, 'view count', fatal=False))
--- a/youtube_dl/extractor/facebook.py
+++ b/youtube_dl/extractor/facebook.py
@ -5,12 +5,14 @@ import re
 import socket
 from .common import InfoExtractor
-from ..utils import (
+from ..compat import (
    compat_http_client,
    compat_str,
    compat_urllib_error,
    compat_urllib_parse,
    compat_urllib_request,
 )
 from ..utils import (
    urlencode_postdata,
    ExtractorError,
    limit_length,
--- a/youtube_dl/extractor/folketinget.py
+++ b/youtube_dl/extractor/folketinget.py
@ -0,0 +1,75 @@
 # -*- coding: utf-8 -*-
 from __future__ import unicode_literals
 from .common import InfoExtractor
 from ..compat import compat_parse_qs
 from ..utils import (
    int_or_none,
    parse_duration,
    parse_iso8601,
    xpath_text,
 )
 class FolketingetIE(InfoExtractor):
    IE_DESC = 'Folketinget (ft.dk; Danish parliament)'
    _VALID_URL = r'https?://(?:www\.)?ft\.dk/webtv/video/[^?#]*?\.(?P<id>[0-9]+)\.aspx'
    _TEST = {
        'url': 'http://www.ft.dk/webtv/video/20141/eru/td.1165642.aspx?as=1#player',
        'info_dict': {
            'id': '1165642',
            'ext': 'mp4',
            'title': 'Åbent samråd i Erhvervsudvalget',
            'description': 'Åbent samråd med erhvervs- og vækstministeren om regeringens politik på teleområdet',
            'view_count': int,
            'width': 768,
            'height': 432,
            'tbr': 928000,
            'timestamp': 1416493800,
            'upload_date': '20141120',
            'duration': 3960,
        },
        'params': {
            'skip_download': 'rtmpdump required',
        }
    }
    def _real_extract(self, url):
        video_id = self._match_id(url)
        webpage = self._download_webpage(url, video_id)
        title = self._og_search_title(webpage)
        description = self._html_search_regex(
            r'(?s)<div class="video-item-agenda"[^>]*>(.*?)<',
            webpage, 'description', fatal=False)
        player_params = compat_parse_qs(self._search_regex(
            r'<embed src="http://ft\.arkena\.tv/flash/ftplayer\.swf\?([^"]+)"',
            webpage, 'player params'))
        xml_url = player_params['xml'][0]
        doc = self._download_xml(xml_url, video_id)
        timestamp = parse_iso8601(xpath_text(doc, './/date'))
        duration = parse_duration(xpath_text(doc, './/duration'))
        width = int_or_none(xpath_text(doc, './/width'))
        height = int_or_none(xpath_text(doc, './/height'))
        view_count = int_or_none(xpath_text(doc, './/views'))
        formats = [{
            'format_id': n.attrib['bitrate'],
            'url': xpath_text(n, './url', fatal=True),
            'tbr': int_or_none(n.attrib['bitrate']),
        } for n in doc.findall('.//streams/stream')]
        self._sort_formats(formats)
        return {
            'id': video_id,
            'title': title,
            'formats': formats,
            'description': description,
            'timestamp': timestamp,
            'width': width,
            'height': height,
            'duration': duration,
            'view_count': view_count,
        }
--- a/youtube_dl/extractor/francetv.py
+++ b/youtube_dl/extractor/francetv.py
@ -26,6 +26,19 @@ class FranceTVBaseInfoExtractor(InfoExtractor):
        if info.get('status') == 'NOK':
            raise ExtractorError(
                '%s returned error: %s' % (self.IE_NAME, info['message']), expected=True)
        allowed_countries = info['videos'][0].get('geoblocage')
        if allowed_countries:
            georestricted = True
            geo_info = self._download_json(
                'http://geo.francetv.fr/ws/edgescape.json', video_id,
                'Downloading geo restriction info')
            country = geo_info['reponse']['geo_info']['country_code']
            if country not in allowed_countries:
                raise ExtractorError(
                    'The video is not available from your location',
                    expected=True)
        else:
            georestricted = False
        formats = []
        for video in info['videos']:
@ -36,6 +49,10 @@ class FranceTVBaseInfoExtractor(InfoExtractor):
                continue
            format_id = video['format']
            if video_url.endswith('.f4m'):
                if georestricted:
                    # See https://github.com/rg3/youtube-dl/issues/3963
                    # m3u8 urls work fine
                    continue
                video_url_parsed = compat_urllib_parse_urlparse(video_url)
                f4m_url = self._download_webpage(
                    'http://hdfauth.francetv.fr/esi/urltokengen2.html?url=%s' % video_url_parsed.path,
--- a/youtube_dl/extractor/freevideo.py
+++ b/youtube_dl/extractor/freevideo.py
@ -0,0 +1,38 @@
 from __future__ import unicode_literals
 from .common import InfoExtractor
 from ..utils import ExtractorError
 class FreeVideoIE(InfoExtractor):
    _VALID_URL = r'^http://www.freevideo.cz/vase-videa/(?P<id>[^.]+)\.html(?:$|[?#])'
    _TEST = {
        'url': 'http://www.freevideo.cz/vase-videa/vysukany-zadecek-22033.html',
        'info_dict': {
            'id': 'vysukany-zadecek-22033',
            'ext': 'mp4',
            "title": "vysukany-zadecek-22033",
            "age_limit": 18,
        },
        'skip': 'Blocked outside .cz',
    }
    def _real_extract(self, url):
        video_id = self._match_id(url)
        webpage, handle = self._download_webpage_handle(url, video_id)
        if '//www.czechav.com/' in handle.geturl():
            raise ExtractorError(
                'Access to freevideo is blocked from your location',
                expected=True)
        video_url = self._search_regex(
            r'\s+url: "(http://[a-z0-9-]+.cdn.freevideo.cz/stream/.*?/video.mp4)"',
            webpage, 'video URL')
        return {
            'id': video_id,
            'url': video_url,
            'title': video_id,
            'age_limit': 18,
        }
--- a/youtube_dl/extractor/funnyordie.py
+++ b/youtube_dl/extractor/funnyordie.py
@ -21,7 +21,6 @@ class FunnyOrDieIE(InfoExtractor):
        },
    }, {
        'url': 'http://www.funnyordie.com/embed/e402820827',
        'md5': '29f4c5e5a61ca39dfd7e8348a75d0aad',
        'info_dict': {
            'id': 'e402820827',
            'ext': 'mp4',
--- a/youtube_dl/extractor/gamekings.py
+++ b/youtube_dl/extractor/gamekings.py
@ -11,7 +11,7 @@ class GamekingsIE(InfoExtractor):
        'url': 'http://www.gamekings.tv/videos/phoenix-wright-ace-attorney-dual-destinies-review/',
        # MD5 is flaky, seems to change regularly
        # 'md5': '2f32b1f7b80fdc5cb616efb4f387f8a3',
-        u'info_dict': {
+        'info_dict': {
            'id': '20130811',
            'ext': 'mp4',
            'title': 'Phoenix Wright: Ace Attorney \u2013 Dual Destinies Review',
--- a/youtube_dl/extractor/gamespot.py
+++ b/youtube_dl/extractor/gamespot.py
@ -8,12 +8,11 @@ from ..utils import (
    compat_urllib_parse,
    compat_urlparse,
    unescapeHTML,
    get_meta_content,
 )
 class GameSpotIE(InfoExtractor):
-    _VALID_URL = r'(?:http://)?(?:www\.)?gamespot\.com/.*-(?P<page_id>\d+)/?'
+    _VALID_URL = r'(?:http://)?(?:www\.)?gamespot\.com/.*-(?P<id>\d+)/?'
    _TEST = {
        'url': 'http://www.gamespot.com/videos/arma-3-community-guide-sitrep-i/2300-6410818/',
        'md5': 'b2a30deaa8654fcccd43713a6b6a4825',
@ -26,10 +25,10 @@ class GameSpotIE(InfoExtractor):
    }
    def _real_extract(self, url):
-        mobj = re.match(self._VALID_URL, url)
+        page_id = self._match_id(url)
        page_id = mobj.group('page_id')
        webpage = self._download_webpage(url, page_id)
-        data_video_json = self._search_regex(r'data-video=["\'](.*?)["\']', webpage, 'data video')
+        data_video_json = self._search_regex(
            r'data-video=["\'](.*?)["\']', webpage, 'data video')
        data_video = json.loads(unescapeHTML(data_video_json))
        # Transform the manifest url to a link to the mp4 files
@ -41,7 +40,8 @@ class GameSpotIE(InfoExtractor):
        http_path = f4m_path[1:].split('/', 1)[1]
        http_template = re.sub(QUALITIES_RE, r'%s', http_path)
        http_template = http_template.replace('.csmil/manifest.f4m', '')
-        http_template = compat_urlparse.urljoin('http://video.gamespotcdn.com/', http_template)
+        http_template = compat_urlparse.urljoin(
            'http://video.gamespotcdn.com/', http_template)
        formats = []
        for q in qualities:
            formats.append({
@ -52,8 +52,9 @@ class GameSpotIE(InfoExtractor):
        return {
            'id': data_video['guid'],
            'display_id': page_id,
            'title': compat_urllib_parse.unquote(data_video['title']),
            'formats': formats,
-            'description': get_meta_content('description', webpage),
+            'description': self._html_search_meta('description', webpage),
            'thumbnail': self._og_search_thumbnail(webpage),
        }
--- a/youtube_dl/extractor/generic.py
+++ b/youtube_dl/extractor/generic.py
@ -7,11 +7,12 @@ import re
 from .common import InfoExtractor
 from .youtube import YoutubeIE
-from ..utils import (
+from ..compat import (
    compat_urllib_parse,
    compat_urlparse,
    compat_xml_parse_error,
-
+)
 from ..utils import (
    determine_ext,
    ExtractorError,
    float_or_none,
@ -99,6 +100,22 @@ class GenericIE(InfoExtractor):
                'uploader': 'Championat',
            },
        },
        {
            # https://github.com/rg3/youtube-dl/issues/3541
            'add_ie': ['Brightcove'],
            'url': 'http://www.kijk.nl/sbs6/leermijvrouwenkennen/videos/jqMiXKAYan2S/aflevering-1',
            'info_dict': {
                'id': '3866516442001',
                'ext': 'mp4',
                'title': 'Leer mij vrouwen kennen: Aflevering 1',
                'description': 'Leer mij vrouwen kennen: Aflevering 1',
                'uploader': 'SBS Broadcasting',
            },
            'skip': 'Restricted to Netherlands',
            'params': {
                'skip_download': True,  # m3u8 download
            },
        },
        # Direct link to a video
        {
            'url': 'http://media.w3.org/2010/05/sintel/trailer.mp4',
@ -417,7 +434,41 @@ class GenericIE(InfoExtractor):
                'title': 'Chet Chat 171 - Oct 29, 2014',
                'upload_date': '20141029',
            }
        },
        # Livestream embed
        {
            'url': 'http://www.esa.int/Our_Activities/Space_Science/Rosetta/Philae_comet_touch-down_webcast',
            'info_dict': {
                'id': '67864563',
                'ext': 'flv',
                'upload_date': '20141112',
                'title': 'Rosetta #CometLanding webcast HL 10',
            }
        },
        # LazyYT
        {
            'url': 'http://discourse.ubuntu.com/t/unity-8-desktop-mode-windows-on-mir/1986',
            'info_dict': {
                'title': 'Unity 8 desktop-mode windows on Mir! - Ubuntu Discourse',
            },
            'playlist_mincount': 2,
        },
        # Direct link with incorrect MIME type
        {
            'url': 'http://ftp.nluug.nl/video/nluug/2014-11-20_nj14/zaal-2/5_Lennart_Poettering_-_Systemd.webm',
            'md5': '4ccbebe5f36706d85221f204d7eb5913',
            'info_dict': {
                'url': 'http://ftp.nluug.nl/video/nluug/2014-11-20_nj14/zaal-2/5_Lennart_Poettering_-_Systemd.webm',
                'id': '5_Lennart_Poettering_-_Systemd',
                'ext': 'webm',
                'title': '5_Lennart_Poettering_-_Systemd',
                'upload_date': '20141120',
            },
            'expected_warnings': [
                'URL could be a direct video link, returning it as such.'
            ]
        }
    ]
    def report_following_redirect(self, new_url):
@ -510,9 +561,9 @@ class GenericIE(InfoExtractor):
            if default_search in ('error', 'fixup_error'):
                raise ExtractorError(
-                    ('%r is not a valid URL. '
+                    '%r is not a valid URL. '
                    'Set --default-search "ytsearch" (or run  youtube-dl "ytsearch:%s" ) to search YouTube'
-                    ) % (url, url), expected=True)
+                    % (url, url), expected=True)
            else:
                if ':' not in default_search:
                    default_search += ':'
@ -559,6 +610,7 @@ class GenericIE(InfoExtractor):
            return {
                'id': video_id,
                'title': os.path.splitext(url_basename(url))[0],
                'direct': True,
                'formats': [{
                    'format_id': m.group('format_id'),
                    'url': url,
@ -570,10 +622,28 @@ class GenericIE(InfoExtractor):
        if not self._downloader.params.get('test', False) and not is_intentional:
            self._downloader.report_warning('Falling back on generic information extractor.')
-        if full_response:
+        if not full_response:
-            webpage = self._webpage_read_content(full_response, url, video_id)
+            full_response = self._request_webpage(url, video_id)
-        else:
+
-            webpage = self._download_webpage(url, video_id)
+        # Maybe it's a direct link to a video?
        # Be careful not to download the whole thing!
        first_bytes = full_response.read(512)
        if not re.match(r'^\s*<', first_bytes.decode('utf-8', 'replace')):
            self._downloader.report_warning(
                'URL could be a direct video link, returning it as such.')
            upload_date = unified_strdate(
                head_response.headers.get('Last-Modified'))
            return {
                'id': video_id,
                'title': os.path.splitext(url_basename(url))[0],
                'direct': True,
                'url': url,
                'upload_date': upload_date,
            }
        webpage = self._webpage_read_content(
            full_response, url, video_id, prefix=first_bytes)
        self.report_extraction(video_id)
        # Is it an RSS feed?
@ -674,6 +744,12 @@ class GenericIE(InfoExtractor):
            return _playlist_from_matches(
                matches, lambda m: unescapeHTML(m[1]))
        # Look for lazyYT YouTube embed
        matches = re.findall(
            r'class="lazyYT" data-youtube-id="([^"]+)"', webpage)
        if matches:
            return _playlist_from_matches(matches, lambda m: unescapeHTML(m))
        # Look for embedded Dailymotion player
        matches = re.findall(
            r'<iframe[^>]+?src=(["\'])(?P<url>(?:https?:)?//(?:www\.)?dailymotion\.com/embed/video/.+?)\1', webpage)
@ -898,6 +974,12 @@ class GenericIE(InfoExtractor):
        if mobj is not None:
            return self.url_result(self._proto_relative_url(mobj.group('url'), scheme='http:'), 'CondeNast')
        mobj = re.search(
            r'<iframe[^>]+src="(?P<url>https?://new\.livestream\.com/[^"]+/player[^"]+)"',
            webpage)
        if mobj is not None:
            return self.url_result(mobj.group('url'), 'Livestream')
        def check_video(vurl):
            vpath = compat_urlparse.urlparse(vurl).path
            vext = determine_ext(vpath)
@ -945,7 +1027,7 @@ class GenericIE(InfoExtractor):
                found = filter_video(re.findall(r'<meta.*?property="og:video".*?content="(.*?)"', webpage))
        if not found:
            # HTML5 video
-            found = re.findall(r'(?s)<video[^<]*(?:>.*?<source[^>]*)?\s+src="([^"]+)"', webpage)
+            found = re.findall(r'(?s)<video[^<]*(?:>.*?<source[^>]*)?\s+src=["\'](.*?)["\']', webpage)
        if not found:
            found = re.search(
                r'(?i)<meta\s+(?=(?:[a-z-]+="[^"]+"\s+)*http-equiv="refresh")'
@ -991,4 +1073,3 @@ class GenericIE(InfoExtractor):
                '_type': 'playlist',
                'entries': entries,
            }
--- a/youtube_dl/extractor/globo.py
+++ b/youtube_dl/extractor/globo.py
@ -5,13 +5,15 @@ import random
 import math
 from .common import InfoExtractor
-from ..utils import (
+from ..compat import (
    ExtractorError,
    float_or_none,
    compat_str,
    compat_chr,
    compat_ord,
 )
 from ..utils import (
    ExtractorError,
    float_or_none,
 )
 class GloboIE(InfoExtractor):
--- a/youtube_dl/extractor/goldenmoustache.py
+++ b/youtube_dl/extractor/goldenmoustache.py
@ -0,0 +1,57 @@
 from __future__ import unicode_literals
 from .common import InfoExtractor
 from ..utils import (
    int_or_none,
 )
 class GoldenMoustacheIE(InfoExtractor):
    _VALID_URL = r'https?://(?:www\.)?goldenmoustache\.com/(?P<display_id>[\w-]+)-(?P<id>\d+)'
    _TESTS = [{
        'url': 'http://www.goldenmoustache.com/suricate-le-poker-3700/',
        'md5': '0f904432fa07da5054d6c8beb5efb51a',
        'info_dict': {
            'id': '3700',
            'ext': 'mp4',
            'title': 'Suricate - Le Poker',
            'description': 'md5:3d1f242f44f8c8cb0a106f1fd08e5dc9',
            'thumbnail': 're:^https?://.*\.jpg$',
            'view_count': int,
        }
    }, {
        'url': 'http://www.goldenmoustache.com/le-lab-tout-effacer-mc-fly-et-carlito-55249/',
        'md5': '27f0c50fb4dd5f01dc9082fc67cd5700',
        'info_dict': {
            'id': '55249',
            'ext': 'mp4',
            'title': 'Le LAB - Tout Effacer (Mc Fly et Carlito)',
            'description': 'md5:9b7fbf11023fb2250bd4b185e3de3b2a',
            'thumbnail': 're:^https?://.*\.(?:png|jpg)$',
            'view_count': int,
        }
    }]
    def _real_extract(self, url):
        video_id = self._match_id(url)
        webpage = self._download_webpage(url, video_id)
        video_url = self._html_search_regex(
            r'data-src-type="mp4" data-src="([^"]+)"', webpage, 'video URL')
        title = self._html_search_regex(
            r'<title>(.*?)(?: - Golden Moustache)?</title>', webpage, 'title')
        thumbnail = self._og_search_thumbnail(webpage)
        description = self._og_search_description(webpage)
        view_count = int_or_none(self._html_search_regex(
            r'<strong>([0-9]+)</strong>\s*VUES</span>',
            webpage, 'view count', fatal=False))
        return {
            'id': video_id,
            'url': video_url,
            'ext': 'mp4',
            'title': title,
            'description': description,
            'thumbnail': thumbnail,
            'view_count': view_count,
        }
--- a/youtube_dl/extractor/gorillavid.py
+++ b/youtube_dl/extractor/gorillavid.py
@ -9,14 +9,15 @@ from ..utils import (
    determine_ext,
    compat_urllib_parse,
    compat_urllib_request,
    int_or_none,
 )
 class GorillaVidIE(InfoExtractor):
-    IE_DESC = 'GorillaVid.in, daclips.in and movpod.in'
+    IE_DESC = 'GorillaVid.in, daclips.in, movpod.in and fastvideo.in'
    _VALID_URL = r'''(?x)
        https?://(?P<host>(?:www\.)?
-            (?:daclips\.in|gorillavid\.in|movpod\.in))/
+            (?:daclips\.in|gorillavid\.in|movpod\.in|fastvideo\.in))/
        (?:embed-)?(?P<id>[0-9a-zA-Z]+)(?:-[0-9]+x[0-9]+\.html)?
    '''
@ -49,6 +50,16 @@ class GorillaVidIE(InfoExtractor):
            'title': 'Micro Pig piglets ready on 16th July 2009-bG0PdrCdxUc',
            'thumbnail': 're:http://.*\.jpg',
        }
    }, {
        # video with countdown timeout
        'url': 'http://fastvideo.in/1qmdn1lmsmbw',
        'md5': '8b87ec3f6564a3108a0e8e66594842ba',
        'info_dict': {
            'id': '1qmdn1lmsmbw',
            'ext': 'mp4',
            'title': 'Man of Steel - Trailer',
            'thumbnail': 're:http://.*\.jpg',
        },
    }, {
        'url': 'http://movpod.in/0wguyyxi1yca',
        'only_matching': True,
@ -71,6 +82,12 @@ class GorillaVidIE(InfoExtractor):
            ''', webpage))
        if fields['op'] == 'download1':
            countdown = int_or_none(self._search_regex(
                r'<span id="countdown_str">(?:[Ww]ait)?\s*<span id="cxc">(\d+)</span>\s*(?:seconds?)?</span>',
                webpage, 'countdown', default=None))
            if countdown:
                self._sleep(countdown, video_id)
            post = compat_urllib_parse.urlencode(fields)
            req = compat_urllib_request.Request(url, post)
@ -78,9 +95,13 @@ class GorillaVidIE(InfoExtractor):
            webpage = self._download_webpage(req, video_id, 'Downloading video page')
-        title = self._search_regex(r'style="z-index: [0-9]+;">([^<]+)</span>', webpage, 'title')
+        title = self._search_regex(
-        video_url = self._search_regex(r'file\s*:\s*\'(http[^\']+)\',', webpage, 'file url')
+            r'style="z-index: [0-9]+;">([^<]+)</span>',
-        thumbnail = self._search_regex(r'image\s*:\s*\'(http[^\']+)\',', webpage, 'thumbnail', fatal=False)
+            webpage, 'title', default=None) or self._og_search_title(webpage)
        video_url = self._search_regex(
            r'file\s*:\s*["\'](http[^"\']+)["\'],', webpage, 'file url')
        thumbnail = self._search_regex(
            r'image\s*:\s*["\'](http[^"\']+)["\'],', webpage, 'thumbnail', fatal=False)
        formats = [{
            'format_id': 'sd',
--- a/youtube_dl/extractor/goshgay.py
+++ b/youtube_dl/extractor/goshgay.py
@ -1,15 +1,11 @@
 # -*- coding: utf-8 -*-
 from __future__ import unicode_literals
 import re
 from .common import InfoExtractor
 from ..utils import (
    compat_urlparse,
    str_to_int,
    ExtractorError,
 )
 import json
 class GoshgayIE(InfoExtractor):
@ -27,36 +23,27 @@ class GoshgayIE(InfoExtractor):
    }
    def _real_extract(self, url):
-        mobj = re.match(self._VALID_URL, url)
+        video_id = self._match_id(url)
        video_id = mobj.group('id')
        webpage = self._download_webpage(url, video_id)
-        title = self._search_regex(r'class="video-title"><h1>(.+?)<', webpage, 'title')
+        title = self._og_search_title(webpage)
        thumbnail = self._og_search_thumbnail(webpage)
        family_friendly = self._html_search_meta(
            'isFamilyFriendly', webpage, default='false')
        config_url = self._search_regex(
            r"'config'\s*:\s*'([^']+)'", webpage, 'config URL')
-        player_config = self._search_regex(
+        config = self._download_xml(
-            r'(?s)jwplayer\("player"\)\.setup\(({.+?})\)', webpage, 'config settings')
+            config_url, video_id, 'Downloading player config XML')
        player_vars = json.loads(player_config.replace("'", '"'))
        width = str_to_int(player_vars.get('width'))
        height = str_to_int(player_vars.get('height'))
        config_uri = player_vars.get('config')
-        if config_uri is None:
+        if config is None:
            raise ExtractorError('Missing config URI')
        node = self._download_xml(config_uri, video_id, 'Downloading player config XML',
                                  errnote='Unable to download XML')
        if node is None:
            raise ExtractorError('Missing config XML')
-        if node.tag != 'config':
+        if config.tag != 'config':
            raise ExtractorError('Missing config attribute')
-        fns = node.findall('file')
+        fns = config.findall('file')
-        imgs = node.findall('image')
+        if len(fns) < 1:
        if len(fns) != 1:
            raise ExtractorError('Missing media URI')
        video_url = fns[0].text
        if len(imgs) < 1:
            thumbnail = None
        else:
            thumbnail = imgs[0].text
        url_comp = compat_urlparse.urlparse(url)
        ref = "%s://%s%s" % (url_comp[0], url_comp[1], url_comp[2])
@ -65,9 +52,7 @@ class GoshgayIE(InfoExtractor):
            'id': video_id,
            'url': video_url,
            'title': title,
            'width': width,
            'height': height,
            'thumbnail': thumbnail,
            'http_referer': ref,
-            'age_limit': 18,
+            'age_limit': 0 if family_friendly == 'true' else 18,
        }
--- a/youtube_dl/extractor/grooveshark.py
+++ b/youtube_dl/extractor/grooveshark.py
@ -8,12 +8,13 @@ import re
 from .common import InfoExtractor
-from ..utils import ExtractorError, compat_urllib_request, compat_html_parser
+from ..compat import (
-
+    compat_html_parser,
 from ..utils import (
    compat_urllib_parse,
    compat_urllib_request,
    compat_urlparse,
 )
 from ..utils import ExtractorError
 class GroovesharkHtmlParser(compat_html_parser.HTMLParser):
--- a/Show More
+++ b/Show More