Update Standards-Version to 3.9.8 (no changes required).

[youtubedl] / youtube_dl / YoutubeDL.py
diff --git a/youtube_dl/YoutubeDL.py b/youtube_dl/YoutubeDL.py

index f4324039c72ec656f67bf536ee7856aae464cc01..e844dc98a5b3915070ffae079395233de7ed04f7 100755 (executable)
--- a/youtube_dl/YoutubeDL.py
+++ b/youtube_dl/YoutubeDL.py
@@ -5,6 +5,7 @@ from __future__ import absolute_import, unicode_literals
  
  import collections
  import contextlib
  
  import collections
  import contextlib
+import copy
  import datetime
  import errno
  import fileinput
  import datetime
  import errno
  import fileinput
@@ -24,9 +25,6 @@ import time
  import tokenize
  import traceback
  
  import tokenize
  import traceback
  
-if os.name == 'nt':
-    import ctypes
-
  from .compat import (
      compat_basestring,
      compat_cookiejar,
  from .compat import (
      compat_basestring,
      compat_cookiejar,
@@ -34,6 +32,7 @@ from .compat import (
      compat_get_terminal_size,
      compat_http_client,
      compat_kwargs,
      compat_get_terminal_size,
      compat_http_client,
      compat_kwargs,
+    compat_os_name,
      compat_str,
      compat_tokenize_tokenize,
      compat_urllib_error,
      compat_str,
      compat_tokenize_tokenize,
      compat_urllib_error,
@@ -41,6 +40,8 @@ from .compat import (
      compat_urllib_request_DataHandler,
  )
  from .utils import (
      compat_urllib_request_DataHandler,
  )
  from .utils import (
+    age_restricted,
+    args_to_str,
      ContentTooShortError,
      date_from_str,
      DateRange,
      ContentTooShortError,
      date_from_str,
      DateRange,
@@ -60,13 +61,17 @@ from .utils import (
      PagedList,
      parse_filesize,
      PerRequestProxyHandler,
      PagedList,
      parse_filesize,
      PerRequestProxyHandler,
-    PostProcessingError,
      platform_name,
      platform_name,
+    PostProcessingError,
      preferredencoding,
      preferredencoding,
+    prepend_extension,
+    register_socks_protocols,
      render_table,
      render_table,
+    replace_extension,
      SameFileError,
      sanitize_filename,
      sanitize_path,
      SameFileError,
      sanitize_filename,
      sanitize_path,
+    sanitize_url,
      sanitized_Request,
      std_headers,
      subtitles_filename,
      sanitized_Request,
      std_headers,
      subtitles_filename,
@@ -77,16 +82,13 @@ from .utils import (
      write_string,
      YoutubeDLCookieProcessor,
      YoutubeDLHandler,
      write_string,
      YoutubeDLCookieProcessor,
      YoutubeDLHandler,
-    prepend_extension,
-    replace_extension,
-    args_to_str,
-    age_restricted,
  )
  from .cache import Cache
  )
  from .cache import Cache
-from .extractor import get_info_extractor, gen_extractors
+from .extractor import get_info_extractor, gen_extractor_classes, _LAZY_LOADER
  from .downloader import get_suitable_downloader
  from .downloader.rtmp import rtmpdump_version
  from .postprocessor import (
  from .downloader import get_suitable_downloader
  from .downloader.rtmp import rtmpdump_version
  from .postprocessor import (
+    FFmpegFixupM3u8PP,
      FFmpegFixupM4aPP,
      FFmpegFixupStretchedPP,
      FFmpegMergerPP,
      FFmpegFixupM4aPP,
      FFmpegFixupStretchedPP,
      FFmpegMergerPP,
@@ -95,6 +97,9 @@ from .postprocessor import (
  )
  from .version import __version__
  
  )
  from .version import __version__
  
+if compat_os_name == 'nt':
+    import ctypes
+
  
  class YoutubeDL(object):
      """YoutubeDL class.
  
  class YoutubeDL(object):
      """YoutubeDL class.
@@ -192,8 +197,8 @@ class YoutubeDL(object):
      prefer_insecure:   Use HTTP instead of HTTPS to retrieve information.
                         At the moment, this is only supported by YouTube.
      proxy:             URL of the proxy server to use
      prefer_insecure:   Use HTTP instead of HTTPS to retrieve information.
                         At the moment, this is only supported by YouTube.
      proxy:             URL of the proxy server to use
-    cn_verification_proxy:  URL of the proxy to use for IP address verification
-                       on Chinese sites. (Experimental)
+    geo_verification_proxy:  URL of the proxy to use for IP address verification
+                       on geo-restricted sites. (Experimental)
      socket_timeout:    Time to wait for unresponsive hosts, in seconds
      bidi_workaround:   Work around buggy terminals without bidirectional text
                         support, using fridibi
      socket_timeout:    Time to wait for unresponsive hosts, in seconds
      bidi_workaround:   Work around buggy terminals without bidirectional text
                         support, using fridibi
@@ -244,7 +249,16 @@ class YoutubeDL(object):
      source_address:    (Experimental) Client-side IP address to bind to.
      call_home:         Boolean, true iff we are allowed to contact the
                         youtube-dl servers for debugging.
      source_address:    (Experimental) Client-side IP address to bind to.
      call_home:         Boolean, true iff we are allowed to contact the
                         youtube-dl servers for debugging.
-    sleep_interval:    Number of seconds to sleep before each download.
+    sleep_interval:    Number of seconds to sleep before each download when
+                       used alone or a lower bound of a range for randomized
+                       sleep before each download (minimum possible number
+                       of seconds to sleep) when used along with
+                       max_sleep_interval.
+    max_sleep_interval:Upper bound of a range for randomized sleep before each
+                       download (maximum possible number of seconds to sleep).
+                       Must only be used along with sleep_interval.
+                       Actual sleep time will be a random float from range
+                       [sleep_interval; max_sleep_interval].
      listformats:       Print an overview of available video formats and exit.
      list_thumbnails:   Print a table of all thumbnails and exit.
      match_filter:      A function that gets called with the info_dict of
      listformats:       Print an overview of available video formats and exit.
      list_thumbnails:   Print a table of all thumbnails and exit.
      match_filter:      A function that gets called with the info_dict of
@@ -257,7 +271,9 @@ class YoutubeDL(object):
      The following options determine which downloader is picked:
      external_downloader: Executable of the external downloader to call.
                         None or unset for standard (built-in) downloader.
      The following options determine which downloader is picked:
      external_downloader: Executable of the external downloader to call.
                         None or unset for standard (built-in) downloader.
-    hls_prefer_native: Use the native HLS downloader instead of ffmpeg/avconv.
+    hls_prefer_native: Use the native HLS downloader instead of ffmpeg/avconv
+                       if True, otherwise use ffmpeg/avconv if False, otherwise
+                       use downloader suggested by extractor if None.
  
      The following parameters are not used by YoutubeDL itself, they are used by
      the downloader (see youtube_dl/downloader/common.py):
  
      The following parameters are not used by YoutubeDL itself, they are used by
      the downloader (see youtube_dl/downloader/common.py):
@@ -298,6 +314,11 @@ class YoutubeDL(object):
          self.params.update(params)
          self.cache = Cache(self)
  
          self.params.update(params)
          self.cache = Cache(self)
  
+        if self.params.get('cn_verification_proxy') is not None:
+            self.report_warning('--cn-verification-proxy is deprecated. Use --geo-verification-proxy instead.')
+            if self.params.get('geo_verification_proxy') is None:
+                self.params['geo_verification_proxy'] = self.params['cn_verification_proxy']
+
          if params.get('bidi_workaround', False):
              try:
                  import pty
          if params.get('bidi_workaround', False):
              try:
                  import pty
@@ -320,7 +341,7 @@ class YoutubeDL(object):
                          ['fribidi', '-c', 'UTF-8'] + width_args, **sp_kwargs)
                  self._output_channel = os.fdopen(master, 'rb')
              except OSError as ose:
                          ['fribidi', '-c', 'UTF-8'] + width_args, **sp_kwargs)
                  self._output_channel = os.fdopen(master, 'rb')
              except OSError as ose:
-                if ose.errno == 2:
+                if ose.errno == errno.ENOENT:
                      self.report_warning('Could not find fribidi executable, ignoring --bidi-workaround . Make sure that  fribidi  is an executable file in one of the directories in your $PATH.')
                  else:
                      raise
                      self.report_warning('Could not find fribidi executable, ignoring --bidi-workaround . Make sure that  fribidi  is an executable file in one of the directories in your $PATH.')
                  else:
                      raise
@@ -356,6 +377,8 @@ class YoutubeDL(object):
          for ph in self.params.get('progress_hooks', []):
              self.add_progress_hook(ph)
  
          for ph in self.params.get('progress_hooks', []):
              self.add_progress_hook(ph)
  
+        register_socks_protocols()
+
      def warn_if_short_id(self, argv):
          # short YouTube ID starting with dash?
          idxs = [
      def warn_if_short_id(self, argv):
          # short YouTube ID starting with dash?
          idxs = [
@@ -375,8 +398,9 @@ class YoutubeDL(object):
      def add_info_extractor(self, ie):
          """Add an InfoExtractor object to the end of the list."""
          self._ies.append(ie)
      def add_info_extractor(self, ie):
          """Add an InfoExtractor object to the end of the list."""
          self._ies.append(ie)
-        self._ies_instances[ie.ie_key()] = ie
-        ie.set_downloader(self)
+        if not isinstance(ie, type):
+            self._ies_instances[ie.ie_key()] = ie
+            ie.set_downloader(self)
  
      def get_info_extractor(self, ie_key):
          """
  
      def get_info_extractor(self, ie_key):
          """
@@ -394,7 +418,7 @@ class YoutubeDL(object):
          """
          Add the InfoExtractors returned by gen_extractors to the end of the list
          """
          """
          Add the InfoExtractors returned by gen_extractors to the end of the list
          """
-        for ie in gen_extractors():
+        for ie in gen_extractor_classes():
              self.add_info_extractor(ie)
  
      def add_post_processor(self, pp):
              self.add_info_extractor(ie)
  
      def add_post_processor(self, pp):
@@ -450,7 +474,7 @@ class YoutubeDL(object):
      def to_console_title(self, message):
          if not self.params.get('consoletitle', False):
              return
      def to_console_title(self, message):
          if not self.params.get('consoletitle', False):
              return
-        if os.name == 'nt' and ctypes.windll.kernel32.GetConsoleWindow():
+        if compat_os_name == 'nt' and ctypes.windll.kernel32.GetConsoleWindow():
              # c_wchar_p() might not be necessary if `message` is
              # already of type unicode()
              ctypes.windll.kernel32.SetConsoleTitleW(ctypes.c_wchar_p(message))
              # c_wchar_p() might not be necessary if `message` is
              # already of type unicode()
              ctypes.windll.kernel32.SetConsoleTitleW(ctypes.c_wchar_p(message))
@@ -521,7 +545,7 @@ class YoutubeDL(object):
          else:
              if self.params.get('no_warnings'):
                  return
          else:
              if self.params.get('no_warnings'):
                  return
-            if not self.params.get('no_color') and self._err_file.isatty() and os.name != 'nt':
+            if not self.params.get('no_color') and self._err_file.isatty() and compat_os_name != 'nt':
                  _msg_header = '\033[0;33mWARNING:\033[0m'
              else:
                  _msg_header = 'WARNING:'
                  _msg_header = '\033[0;33mWARNING:\033[0m'
              else:
                  _msg_header = 'WARNING:'
@@ -533,7 +557,7 @@ class YoutubeDL(object):
          Do the same as trouble, but prefixes the message with 'ERROR:', colored
          in red if stderr is a tty file.
          '''
          Do the same as trouble, but prefixes the message with 'ERROR:', colored
          in red if stderr is a tty file.
          '''
-        if not self.params.get('no_color') and self._err_file.isatty() and os.name != 'nt':
+        if not self.params.get('no_color') and self._err_file.isatty() and compat_os_name != 'nt':
              _msg_header = '\033[0;31mERROR:\033[0m'
          else:
              _msg_header = 'ERROR:'
              _msg_header = '\033[0;31mERROR:\033[0m'
          else:
              _msg_header = 'ERROR:'
@@ -566,7 +590,7 @@ class YoutubeDL(object):
                  elif template_dict.get('height'):
                      template_dict['resolution'] = '%sp' % template_dict['height']
                  elif template_dict.get('width'):
                  elif template_dict.get('height'):
                      template_dict['resolution'] = '%sp' % template_dict['height']
                  elif template_dict.get('width'):
-                    template_dict['resolution'] = '?x%d' % template_dict['width']
+                    template_dict['resolution'] = '%dx?' % template_dict['width']
  
              sanitize = lambda k, v: sanitize_filename(
                  compat_str(v),
  
              sanitize = lambda k, v: sanitize_filename(
                  compat_str(v),
@@ -574,7 +598,7 @@ class YoutubeDL(object):
                  is_id=(k == 'id'))
              template_dict = dict((k, sanitize(k, v))
                                   for k, v in template_dict.items()
                  is_id=(k == 'id'))
              template_dict = dict((k, sanitize(k, v))
                                   for k, v in template_dict.items()
-                                 if v is not None)
+                                 if v is not None and not isinstance(v, (list, tuple, dict)))
              template_dict = collections.defaultdict(lambda: 'NA', template_dict)
  
              outtmpl = self.params.get('outtmpl', DEFAULT_OUTTMPL)
              template_dict = collections.defaultdict(lambda: 'NA', template_dict)
  
              outtmpl = self.params.get('outtmpl', DEFAULT_OUTTMPL)
@@ -658,6 +682,7 @@ class YoutubeDL(object):
              if not ie.suitable(url):
                  continue
  
              if not ie.suitable(url):
                  continue
  
+            ie = self.get_info_extractor(ie.ie_key())
              if not ie.working():
                  self.report_warning('The program functionality for this site has been marked as broken, '
                                      'and will probably not work.')
              if not ie.working():
                  self.report_warning('The program functionality for this site has been marked as broken, '
                                      'and will probably not work.')
@@ -710,6 +735,7 @@ class YoutubeDL(object):
          result_type = ie_result.get('_type', 'video')
  
          if result_type in ('url', 'url_transparent'):
          result_type = ie_result.get('_type', 'video')
  
          if result_type in ('url', 'url_transparent'):
+            ie_result['url'] = sanitize_url(ie_result['url'])
              extract_flat = self.params.get('extract_flat', False)
              if ((extract_flat == 'in_playlist' and 'playlist' in extra_info) or
                      extract_flat is True):
              extract_flat = self.params.get('extract_flat', False)
              if ((extract_flat == 'in_playlist' and 'playlist' in extra_info) or
                      extract_flat is True):
@@ -903,7 +929,7 @@ class YoutubeDL(object):
                  '*=': lambda attr, value: value in attr,
              }
              str_operator_rex = re.compile(r'''(?x)
                  '*=': lambda attr, value: value in attr,
              }
              str_operator_rex = re.compile(r'''(?x)
-                \s*(?P<key>ext|acodec|vcodec|container|protocol)
+                \s*(?P<key>ext|acodec|vcodec|container|protocol|format_id)
                  \s*(?P<op>%s)(?P<none_inclusive>\s*\?)?
                  \s*(?P<value>[a-zA-Z0-9._-]+)
                  \s*$
                  \s*(?P<op>%s)(?P<none_inclusive>\s*\?)?
                  \s*(?P<value>[a-zA-Z0-9._-]+)
                  \s*$
@@ -1035,9 +1061,9 @@ class YoutubeDL(object):
              if isinstance(selector, list):
                  fs = [_build_selector_function(s) for s in selector]
  
              if isinstance(selector, list):
                  fs = [_build_selector_function(s) for s in selector]
  
-                def selector_function(formats):
+                def selector_function(ctx):
                      for f in fs:
                      for f in fs:
-                        for format in f(formats):
+                        for format in f(ctx):
                              yield format
                  return selector_function
              elif selector.type == GROUP:
                              yield format
                  return selector_function
              elif selector.type == GROUP:
@@ -1045,17 +1071,17 @@ class YoutubeDL(object):
              elif selector.type == PICKFIRST:
                  fs = [_build_selector_function(s) for s in selector.selector]
  
              elif selector.type == PICKFIRST:
                  fs = [_build_selector_function(s) for s in selector.selector]
  
-                def selector_function(formats):
+                def selector_function(ctx):
                      for f in fs:
                      for f in fs:
-                        picked_formats = list(f(formats))
+                        picked_formats = list(f(ctx))
                          if picked_formats:
                              return picked_formats
                      return []
              elif selector.type == SINGLE:
                  format_spec = selector.selector
  
                          if picked_formats:
                              return picked_formats
                      return []
              elif selector.type == SINGLE:
                  format_spec = selector.selector
  
-                def selector_function(formats):
-                    formats = list(formats)
+                def selector_function(ctx):
+                    formats = list(ctx['formats'])
                      if not formats:
                          return
                      if format_spec == 'all':
                      if not formats:
                          return
                      if format_spec == 'all':
@@ -1068,9 +1094,10 @@ class YoutubeDL(object):
                              if f.get('vcodec') != 'none' and f.get('acodec') != 'none']
                          if audiovideo_formats:
                              yield audiovideo_formats[format_idx]
                              if f.get('vcodec') != 'none' and f.get('acodec') != 'none']
                          if audiovideo_formats:
                              yield audiovideo_formats[format_idx]
-                        # for audio only (soundcloud) or video only (imgur) urls, select the best/worst audio format
-                        elif (all(f.get('acodec') != 'none' for f in formats) or
-                              all(f.get('vcodec') != 'none' for f in formats)):
+                        # for extractors with incomplete formats (audio only (soundcloud)
+                        # or video only (imgur)) we will fallback to best/worst
+                        # {video,audio}-only format
+                        elif ctx['incomplete_formats']:
                              yield formats[format_idx]
                      elif format_spec == 'bestaudio':
                          audio_formats = [
                              yield formats[format_idx]
                      elif format_spec == 'bestaudio':
                          audio_formats = [
@@ -1144,17 +1171,18 @@ class YoutubeDL(object):
                      }
                  video_selector, audio_selector = map(_build_selector_function, selector.selector)
  
                      }
                  video_selector, audio_selector = map(_build_selector_function, selector.selector)
  
-                def selector_function(formats):
-                    formats = list(formats)
-                    for pair in itertools.product(video_selector(formats), audio_selector(formats)):
+                def selector_function(ctx):
+                    for pair in itertools.product(
+                            video_selector(copy.deepcopy(ctx)), audio_selector(copy.deepcopy(ctx))):
                          yield _merge(pair)
  
              filters = [self._build_format_filter(f) for f in selector.filters]
  
                          yield _merge(pair)
  
              filters = [self._build_format_filter(f) for f in selector.filters]
  
-            def final_selector(formats):
+            def final_selector(ctx):
+                ctx_copy = copy.deepcopy(ctx)
                  for _filter in filters:
                  for _filter in filters:
-                    formats = list(filter(_filter, formats))
-                return selector_function(formats)
+                    ctx_copy['formats'] = list(filter(_filter, ctx_copy['formats']))
+                return selector_function(ctx_copy)
              return final_selector
  
          stream = io.BytesIO(format_spec.encode('utf-8'))
              return final_selector
  
          stream = io.BytesIO(format_spec.encode('utf-8'))
@@ -1212,6 +1240,10 @@ class YoutubeDL(object):
          if 'title' not in info_dict:
              raise ExtractorError('Missing "title" field in extractor result')
  
          if 'title' not in info_dict:
              raise ExtractorError('Missing "title" field in extractor result')
  
+        if not isinstance(info_dict['id'], compat_str):
+            self.report_warning('"id" field is not a string - forcing string conversion')
+            info_dict['id'] = compat_str(info_dict['id'])
+
          if 'playlist' not in info_dict:
              # It isn't part of a playlist
              info_dict['playlist'] = None
          if 'playlist' not in info_dict:
              # It isn't part of a playlist
              info_dict['playlist'] = None
@@ -1227,12 +1259,20 @@ class YoutubeDL(object):
                  t.get('preference'), t.get('width'), t.get('height'),
                  t.get('id'), t.get('url')))
              for i, t in enumerate(thumbnails):
                  t.get('preference'), t.get('width'), t.get('height'),
                  t.get('id'), t.get('url')))
              for i, t in enumerate(thumbnails):
+                t['url'] = sanitize_url(t['url'])
                  if t.get('width') and t.get('height'):
                      t['resolution'] = '%dx%d' % (t['width'], t['height'])
                  if t.get('id') is None:
                      t['id'] = '%d' % i
  
                  if t.get('width') and t.get('height'):
                      t['resolution'] = '%dx%d' % (t['width'], t['height'])
                  if t.get('id') is None:
                      t['id'] = '%d' % i
  
-        if thumbnails and 'thumbnail' not in info_dict:
+        if self.params.get('list_thumbnails'):
+            self.list_thumbnails(info_dict)
+            return
+
+        thumbnail = info_dict.get('thumbnail')
+        if thumbnail:
+            info_dict['thumbnail'] = sanitize_url(thumbnail)
+        elif thumbnails:
              info_dict['thumbnail'] = thumbnails[-1]['url']
  
          if 'display_id' not in info_dict and 'id' in info_dict:
              info_dict['thumbnail'] = thumbnails[-1]['url']
  
          if 'display_id' not in info_dict and 'id' in info_dict:
@@ -1257,6 +1297,8 @@ class YoutubeDL(object):
          if subtitles:
              for _, subtitle in subtitles.items():
                  for subtitle_format in subtitle:
          if subtitles:
              for _, subtitle in subtitles.items():
                  for subtitle_format in subtitle:
+                    if subtitle_format.get('url'):
+                        subtitle_format['url'] = sanitize_url(subtitle_format['url'])
                      if 'ext' not in subtitle_format:
                          subtitle_format['ext'] = determine_ext(subtitle_format['url']).lower()
  
                      if 'ext' not in subtitle_format:
                          subtitle_format['ext'] = determine_ext(subtitle_format['url']).lower()
  
@@ -1286,6 +1328,8 @@ class YoutubeDL(object):
              if 'url' not in format:
                  raise ExtractorError('Missing "url" key in result (index %d)' % i)
  
              if 'url' not in format:
                  raise ExtractorError('Missing "url" key in result (index %d)' % i)
  
+            format['url'] = sanitize_url(format['url'])
+
              if format.get('format_id') is None:
                  format['format_id'] = compat_str(i)
              else:
              if format.get('format_id') is None:
                  format['format_id'] = compat_str(i)
              else:
@@ -1333,9 +1377,6 @@ class YoutubeDL(object):
          if self.params.get('listformats'):
              self.list_formats(info_dict)
              return
          if self.params.get('listformats'):
              self.list_formats(info_dict)
              return
-        if self.params.get('list_thumbnails'):
-            self.list_thumbnails(info_dict)
-            return
  
          req_format = self.params.get('format')
          if req_format is None:
  
          req_format = self.params.get('format')
          if req_format is None:
@@ -1348,7 +1389,34 @@ class YoutubeDL(object):
              req_format_list.append('best')
              req_format = '/'.join(req_format_list)
          format_selector = self.build_format_selector(req_format)
              req_format_list.append('best')
              req_format = '/'.join(req_format_list)
          format_selector = self.build_format_selector(req_format)
-        formats_to_download = list(format_selector(formats))
+
+        # While in format selection we may need to have an access to the original
+        # format set in order to calculate some metrics or do some processing.
+        # For now we need to be able to guess whether original formats provided
+        # by extractor are incomplete or not (i.e. whether extractor provides only
+        # video-only or audio-only formats) for proper formats selection for
+        # extractors with such incomplete formats (see
+        # https://github.com/rg3/youtube-dl/pull/5556).
+        # Since formats may be filtered during format selection and may not match
+        # the original formats the results may be incorrect. Thus original formats
+        # or pre-calculated metrics should be passed to format selection routines
+        # as well.
+        # We will pass a context object containing all necessary additional data
+        # instead of just formats.
+        # This fixes incorrect format selection issue (see
+        # https://github.com/rg3/youtube-dl/issues/10083).
+        incomplete_formats = (
+            # All formats are video-only or
+            all(f.get('vcodec') != 'none' and f.get('acodec') == 'none' for f in formats) or
+            # all formats are audio-only
+            all(f.get('vcodec') == 'none' and f.get('acodec') != 'none' for f in formats))
+
+        ctx = {
+            'formats': formats,
+            'incomplete_formats': incomplete_formats,
+        }
+
+        formats_to_download = list(format_selector(ctx))
          if not formats_to_download:
              raise ExtractorError('requested format not available',
                                   expected=True)
          if not formats_to_download:
              raise ExtractorError('requested format not available',
                                   expected=True)
@@ -1535,7 +1603,9 @@ class YoutubeDL(object):
                          self.to_screen('[info] Video subtitle %s.%s is already_present' % (sub_lang, sub_format))
                      else:
                          self.to_screen('[info] Writing video subtitles to: ' + sub_filename)
                          self.to_screen('[info] Video subtitle %s.%s is already_present' % (sub_lang, sub_format))
                      else:
                          self.to_screen('[info] Writing video subtitles to: ' + sub_filename)
-                        with io.open(encodeFilename(sub_filename), 'w', encoding='utf-8') as subfile:
+                        # Use newline='' to prevent conversion of newline characters
+                        # See https://github.com/rg3/youtube-dl/issues/10268
+                        with io.open(encodeFilename(sub_filename), 'w', encoding='utf-8', newline='') as subfile:
                              subfile.write(sub_data)
                  except (OSError, IOError):
                      self.report_error('Cannot write subtitles file ' + sub_filename)
                              subfile.write(sub_data)
                  except (OSError, IOError):
                      self.report_error('Cannot write subtitles file ' + sub_filename)
@@ -1623,7 +1693,7 @@ class YoutubeDL(object):
                      # Just a single file
                      success = dl(filename, info_dict)
              except (compat_urllib_error.URLError, compat_http_client.HTTPException, socket.error) as err:
                      # Just a single file
                      success = dl(filename, info_dict)
              except (compat_urllib_error.URLError, compat_http_client.HTTPException, socket.error) as err:
-                self.report_error('unable to download video data: %s' % str(err))
+                self.report_error('unable to download video data: %s' % error_to_compat_str(err))
                  return
              except (OSError, IOError) as err:
                  raise UnavailableVideoError(err)
                  return
              except (OSError, IOError) as err:
                  raise UnavailableVideoError(err)
@@ -1631,12 +1701,14 @@ class YoutubeDL(object):
                  self.report_error('content too short (expected %s bytes and served %s)' % (err.expected, err.downloaded))
                  return
  
                  self.report_error('content too short (expected %s bytes and served %s)' % (err.expected, err.downloaded))
                  return
  
-            if success:
+            if success and filename != '-':
                  # Fixup content
                  fixup_policy = self.params.get('fixup')
                  if fixup_policy is None:
                      fixup_policy = 'detect_or_warn'
  
                  # Fixup content
                  fixup_policy = self.params.get('fixup')
                  if fixup_policy is None:
                      fixup_policy = 'detect_or_warn'
  
+                INSTALL_FFMPEG_MESSAGE = 'Install ffmpeg or avconv to fix this automatically.'
+
                  stretched_ratio = info_dict.get('stretched_ratio')
                  if stretched_ratio is not None and stretched_ratio != 1:
                      if fixup_policy == 'warn':
                  stretched_ratio = info_dict.get('stretched_ratio')
                  if stretched_ratio is not None and stretched_ratio != 1:
                      if fixup_policy == 'warn':
@@ -1649,15 +1721,18 @@ class YoutubeDL(object):
                              info_dict['__postprocessors'].append(stretched_pp)
                          else:
                              self.report_warning(
                              info_dict['__postprocessors'].append(stretched_pp)
                          else:
                              self.report_warning(
-                                '%s: Non-uniform pixel ratio (%s). Install ffmpeg or avconv to fix this automatically.' % (
-                                    info_dict['id'], stretched_ratio))
+                                '%s: Non-uniform pixel ratio (%s). %s'
+                                % (info_dict['id'], stretched_ratio, INSTALL_FFMPEG_MESSAGE))
                      else:
                          assert fixup_policy in ('ignore', 'never')
  
                      else:
                          assert fixup_policy in ('ignore', 'never')
  
-                if info_dict.get('requested_formats') is None and info_dict.get('container') == 'm4a_dash':
+                if (info_dict.get('requested_formats') is None and
+                        info_dict.get('container') == 'm4a_dash'):
                      if fixup_policy == 'warn':
                      if fixup_policy == 'warn':
-                        self.report_warning('%s: writing DASH m4a. Only some players support this container.' % (
-                            info_dict['id']))
+                        self.report_warning(
+                            '%s: writing DASH m4a. '
+                            'Only some players support this container.'
+                            % info_dict['id'])
                      elif fixup_policy == 'detect_or_warn':
                          fixup_pp = FFmpegFixupM4aPP(self)
                          if fixup_pp.available:
                      elif fixup_policy == 'detect_or_warn':
                          fixup_pp = FFmpegFixupM4aPP(self)
                          if fixup_pp.available:
@@ -1665,8 +1740,27 @@ class YoutubeDL(object):
                              info_dict['__postprocessors'].append(fixup_pp)
                          else:
                              self.report_warning(
                              info_dict['__postprocessors'].append(fixup_pp)
                          else:
                              self.report_warning(
-                                '%s: writing DASH m4a. Only some players support this container. Install ffmpeg or avconv to fix this automatically.' % (
-                                    info_dict['id']))
+                                '%s: writing DASH m4a. '
+                                'Only some players support this container. %s'
+                                % (info_dict['id'], INSTALL_FFMPEG_MESSAGE))
+                    else:
+                        assert fixup_policy in ('ignore', 'never')
+
+                if (info_dict.get('protocol') == 'm3u8_native' or
+                        info_dict.get('protocol') == 'm3u8' and
+                        self.params.get('hls_prefer_native')):
+                    if fixup_policy == 'warn':
+                        self.report_warning('%s: malformated aac bitstream.' % (
+                            info_dict['id']))
+                    elif fixup_policy == 'detect_or_warn':
+                        fixup_pp = FFmpegFixupM3u8PP(self)
+                        if fixup_pp.available:
+                            info_dict.setdefault('__postprocessors', [])
+                            info_dict['__postprocessors'].append(fixup_pp)
+                        else:
+                            self.report_warning(
+                                '%s: malformated aac bitstream. %s'
+                                % (info_dict['id'], INSTALL_FFMPEG_MESSAGE))
                      else:
                          assert fixup_policy in ('ignore', 'never')
  
                      else:
                          assert fixup_policy in ('ignore', 'never')
  
@@ -1809,7 +1903,7 @@ class YoutubeDL(object):
          if fdict.get('language'):
              if res:
                  res += ' '
          if fdict.get('language'):
              if res:
                  res += ' '
-            res += '[%s]' % fdict['language']
+            res += '[%s] ' % fdict['language']
          if fdict.get('format_note') is not None:
              res += fdict['format_note'] + ' '
          if fdict.get('tbr') is not None:
          if fdict.get('format_note') is not None:
              res += fdict['format_note'] + ' '
          if fdict.get('tbr') is not None:
@@ -1830,7 +1924,9 @@ class YoutubeDL(object):
          if fdict.get('vbr') is not None:
              res += '%4dk' % fdict['vbr']
          if fdict.get('fps') is not None:
          if fdict.get('vbr') is not None:
              res += '%4dk' % fdict['vbr']
          if fdict.get('fps') is not None:
-            res += ', %sfps' % fdict['fps']
+            if res:
+                res += ', '
+            res += '%sfps' % fdict['fps']
          if fdict.get('acodec') is not None:
              if res:
                  res += ', '
          if fdict.get('acodec') is not None:
              if res:
                  res += ', '
@@ -1873,13 +1969,8 @@ class YoutubeDL(object):
      def list_thumbnails(self, info_dict):
          thumbnails = info_dict.get('thumbnails')
          if not thumbnails:
      def list_thumbnails(self, info_dict):
          thumbnails = info_dict.get('thumbnails')
          if not thumbnails:
-            tn_url = info_dict.get('thumbnail')
-            if tn_url:
-                thumbnails = [{'id': '0', 'url': tn_url}]
-            else:
-                self.to_screen(
-                    '[info] No thumbnails present for %s' % info_dict['id'])
-                return
+            self.to_screen('[info] No thumbnails present for %s' % info_dict['id'])
+            return
  
          self.to_screen(
              '[info] Thumbnails for %s:' % info_dict['id'])
  
          self.to_screen(
              '[info] Thumbnails for %s:' % info_dict['id'])
@@ -1924,6 +2015,8 @@ class YoutubeDL(object):
          write_string(encoding_str, encoding=None)
  
          self._write_string('[debug] youtube-dl version ' + __version__ + '\n')
          write_string(encoding_str, encoding=None)
  
          self._write_string('[debug] youtube-dl version ' + __version__ + '\n')
+        if _LAZY_LOADER:
+            self._write_string('[debug] Lazy loading extractors enabled' + '\n')
          try:
              sp = subprocess.Popen(
                  ['git', 'rev-parse', '--short', 'HEAD'],
          try:
              sp = subprocess.Popen(
                  ['git', 'rev-parse', '--short', 'HEAD'],
@@ -1979,6 +2072,7 @@ class YoutubeDL(object):
          if opts_cookiefile is None:
              self.cookiejar = compat_cookiejar.CookieJar()
          else:
          if opts_cookiefile is None:
              self.cookiejar = compat_cookiejar.CookieJar()
          else:
+            opts_cookiefile = compat_expanduser(opts_cookiefile)
              self.cookiejar = compat_cookiejar.MozillaCookieJar(
                  opts_cookiefile)
              if os.access(opts_cookiefile, os.R_OK):
              self.cookiejar = compat_cookiejar.MozillaCookieJar(
                  opts_cookiefile)
              if os.access(opts_cookiefile, os.R_OK):