Merge tag 'upstream/2014.12.01'

[youtubedl] / youtube_dl / extractor / bliptv.py
diff --git a/youtube_dl/extractor/bliptv.py b/youtube_dl/extractor/bliptv.py

index acfc4ad736d9deecf1ed9cdadd4063e2fc8e7243..da47f27bdd6702d3927f3fde72fc0ebe064df53a 100644 (file)
--- a/youtube_dl/extractor/bliptv.py
+++ b/youtube_dl/extractor/bliptv.py
@@ -15,7 +15,7 @@ from ..utils import (
  
  
  class BlipTVIE(SubtitlesInfoExtractor):
  
  
  class BlipTVIE(SubtitlesInfoExtractor):
-    _VALID_URL = r'https?://(?:\w+\.)?blip\.tv/(?:(?:.+-|rss/flash/)(?P<id>\d+)|((?:play/|api\.swf#)(?P<lookup_id>[\da-zA-Z+]+)))'
+    _VALID_URL = r'https?://(?:\w+\.)?blip\.tv/(?:(?:.+-|rss/flash/)(?P<id>\d+)|((?:play/|api\.swf#)(?P<lookup_id>[\da-zA-Z+_]+)))'
  
      _TESTS = [
          {
  
      _TESTS = [
          {
@@ -49,6 +49,35 @@ class BlipTVIE(SubtitlesInfoExtractor):
                  'uploader_id': '792887',
                  'duration': 279,
              }
                  'uploader_id': '792887',
                  'duration': 279,
              }
+        },
+        {
+            # https://bugzilla.redhat.com/show_bug.cgi?id=967465
+            'url': 'http://a.blip.tv/api.swf#h6Uag5KbVwI',
+            'md5': '314e87b1ebe7a48fcbfdd51b791ce5a6',
+            'info_dict': {
+                'id': '6573122',
+                'ext': 'mov',
+                'upload_date': '20130520',
+                'description': 'Two hapless space marines argue over what to do when they realize they have an astronomically huge problem on their hands.',
+                'title': 'Red vs. Blue Season 11 Trailer',
+                'timestamp': 1369029609,
+                'uploader': 'redvsblue',
+                'uploader_id': '792887',
+            }
+        },
+        {
+            'url': 'http://blip.tv/play/gbk766dkj4Yn',
+            'md5': 'fe0a33f022d49399a241e84a8ea8b8e3',
+            'info_dict': {
+                'id': '1749452',
+                'ext': 'mp4',
+                'upload_date': '20090208',
+                'description': 'Witness the first appearance of the Nostalgia Critic character, as Doug reviews the movie Transformers.',
+                'title': 'Nostalgia Critic: Transformers',
+                'timestamp': 1234068723,
+                'uploader': 'NostalgiaCritic',
+                'uploader_id': '246467',
+            }
          }
      ]
  
          }
      ]
  
@@ -56,13 +85,16 @@ class BlipTVIE(SubtitlesInfoExtractor):
          mobj = re.match(self._VALID_URL, url)
          lookup_id = mobj.group('lookup_id')
  
          mobj = re.match(self._VALID_URL, url)
          lookup_id = mobj.group('lookup_id')
  
-        # See https://github.com/rg3/youtube-dl/issues/857
+        # See https://github.com/rg3/youtube-dl/issues/857 and
+        # https://github.com/rg3/youtube-dl/issues/4197
          if lookup_id:
          if lookup_id:
-            info_page = self._download_webpage(
-                'http://blip.tv/play/%s.x?p=1' % lookup_id, lookup_id, 'Resolving lookup id')
-            video_id = self._search_regex(r'data-episode-id="([0-9]+)', info_page, 'video_id')
-        else:
-            video_id = mobj.group('id')
+            urlh = self._request_webpage(
+                'http://blip.tv/play/%s' % lookup_id, lookup_id, 'Resolving lookup id')
+            url = compat_urlparse.urlparse(urlh.geturl())
+            qs = compat_urlparse.parse_qs(url.query)
+            mobj = re.match(self._VALID_URL, qs['file'][0])
+
+        video_id = mobj.group('id')
  
          rss = self._download_xml('http://blip.tv/rss/flash/%s' % video_id, video_id, 'Downloading video RSS')
  
  
          rss = self._download_xml('http://blip.tv/rss/flash/%s' % video_id, video_id, 'Downloading video RSS')
  
@@ -98,7 +130,7 @@ class BlipTVIE(SubtitlesInfoExtractor):
              msg = self._download_webpage(
                  url + '?showplayer=20140425131715&referrer=http://blip.tv&mask=7&skin=flashvars&view=url',
                  video_id, 'Resolving URL for %s' % role)
              msg = self._download_webpage(
                  url + '?showplayer=20140425131715&referrer=http://blip.tv&mask=7&skin=flashvars&view=url',
                  video_id, 'Resolving URL for %s' % role)
-            real_url = compat_urlparse.parse_qs(msg)['message'][0]
+            real_url = compat_urlparse.parse_qs(msg.strip())['message'][0]
  
              media_type = media_content.get('type')
              if media_type == 'text/srt' or url.endswith('.srt'):
  
              media_type = media_content.get('type')
              if media_type == 'text/srt' or url.endswith('.srt'):
@@ -150,9 +182,17 @@ class BlipTVIE(SubtitlesInfoExtractor):
  
  
  class BlipTVUserIE(InfoExtractor):
  
  
  class BlipTVUserIE(InfoExtractor):
-    _VALID_URL = r'(?:(?:(?:https?://)?(?:\w+\.)?blip\.tv/)|bliptvuser:)([^/]+)/*$'
+    _VALID_URL = r'(?:(?:https?://(?:\w+\.)?blip\.tv/)|bliptvuser:)(?!api\.swf)([^/]+)/*$'
      _PAGE_SIZE = 12
      IE_NAME = 'blip.tv:user'
      _PAGE_SIZE = 12
      IE_NAME = 'blip.tv:user'
+    _TEST = {
+        'url': 'http://blip.tv/actone',
+        'info_dict': {
+            'id': 'actone',
+            'title': 'Act One: The Series',
+        },
+        'playlist_count': 5,
+    }
  
      def _real_extract(self, url):
          mobj = re.match(self._VALID_URL, url)
  
      def _real_extract(self, url):
          mobj = re.match(self._VALID_URL, url)
@@ -163,6 +203,7 @@ class BlipTVUserIE(InfoExtractor):
          page = self._download_webpage(url, username, 'Downloading user page')
          mobj = re.search(r'data-users-id="([^"]+)"', page)
          page_base = page_base % mobj.group(1)
          page = self._download_webpage(url, username, 'Downloading user page')
          mobj = re.search(r'data-users-id="([^"]+)"', page)
          page_base = page_base % mobj.group(1)
+        title = self._og_search_title(page)
  
          # Download video ids using BlipTV Ajax calls. Result size per
          # query is limited (currently to 12 videos) so we need to query
  
          # Download video ids using BlipTV Ajax calls. Result size per
          # query is limited (currently to 12 videos) so we need to query
@@ -199,4 +240,5 @@ class BlipTVUserIE(InfoExtractor):
  
          urls = ['http://blip.tv/%s' % video_id for video_id in video_ids]
          url_entries = [self.url_result(vurl, 'BlipTV') for vurl in urls]
  
          urls = ['http://blip.tv/%s' % video_id for video_id in video_ids]
          url_entries = [self.url_result(vurl, 'BlipTV') for vurl in urls]
-        return [self.playlist_result(url_entries, playlist_title=username)]
+        return self.playlist_result(
+            url_entries, playlist_title=title, playlist_id=username)