# coding: utf-8
-from __future__ import unicode_literals, division
+from __future__ import unicode_literals
import re
from .common import InfoExtractor
-from ..compat import (
- compat_chr,
- compat_ord,
-)
+from ..compat import compat_chr
from ..utils import (
determine_ext,
ExtractorError,
)
-from ..jsinterp import (
- JSInterpreter,
- _NAME_RE
-)
class OpenloadIE(InfoExtractor):
- _VALID_URL = r'https?://openload\.(?:co|io)/(?:f|embed)/(?P<id>[a-zA-Z0-9-_]+)'
+ _VALID_URL = r'https?://(?:openload\.(?:co|io)|oload\.tv)/(?:f|embed)/(?P<id>[a-zA-Z0-9-_]+)'
_TESTS = [{
'url': 'https://openload.co/f/kUEfGclsU9o',
'id': 'kUEfGclsU9o',
'ext': 'mp4',
'title': 'skyrim_no-audio_1080.mp4',
- 'thumbnail': 're:^https?://.*\.jpg$',
+ 'thumbnail': r're:^https?://.*\.jpg$',
},
}, {
'url': 'https://openload.co/embed/rjC09fkPLYs',
'id': 'rjC09fkPLYs',
'ext': 'mp4',
'title': 'movie.mp4',
- 'thumbnail': 're:^https?://.*\.jpg$',
+ 'thumbnail': r're:^https?://.*\.jpg$',
'subtitles': {
'en': [{
'ext': 'vtt',
# for title and ext
'url': 'https://openload.co/embed/Sxz5sADo82g/',
'only_matching': True,
+ }, {
+ 'url': 'https://oload.tv/embed/KnG-kKZdcfY/',
+ 'only_matching': True,
}]
- def openload_decode(self, txt):
- symbol_dict = {
- '(゚Д゚) [゚Θ゚]': '_',
- '(゚Д゚) [゚ω゚ノ]': 'a',
- '(゚Д゚) [゚Θ゚ノ]': 'b',
- '(゚Д゚) [\'c\']': 'c',
- '(゚Д゚) [゚ー゚ノ]': 'd',
- '(゚Д゚) [゚Д゚ノ]': 'e',
- '(゚Д゚) [1]': 'f',
- '(゚Д゚) [\'o\']': 'o',
- '(o゚ー゚o)': 'u',
- '(゚Д゚) [\'c\']': 'c',
- '((゚ー゚) + (o^_^o))': '7',
- '((o^_^o) +(o^_^o) +(c^_^o))': '6',
- '((゚ー゚) + (゚Θ゚))': '5',
- '(-~3)': '4',
- '(-~-~1)': '3',
- '(-~1)': '2',
- '(-~0)': '1',
- '((c^_^o)-(c^_^o))': '0',
- }
- delim = '(゚Д゚)[゚ε゚]+'
- end_token = '(゚Д゚)[゚o゚]'
- symbols = '|'.join(map(re.escape, symbol_dict.keys()))
- txt = re.sub('(%s)\+\s?' % symbols, lambda m: symbol_dict[m.group(1)], txt)
- ret = ''
- for aacode in re.findall(r'{0}\+\s?{1}(.*?){0}'.format(re.escape(end_token), re.escape(delim)), txt):
- for aachar in aacode.split(delim):
- if aachar.isdigit():
- ret += compat_chr(int(aachar, 8))
- else:
- m = re.match(r'^u([\da-f]{4})$', aachar)
- if m:
- ret += compat_chr(int(m.group(1), 16))
- else:
- self.report_warning("Cannot decode: %s" % aachar)
- return ret
+ @staticmethod
+ def _extract_urls(webpage):
+ return re.findall(
+ r'<iframe[^>]+src=["\']((?:https?://)?(?:openload\.(?:co|io)|oload\.tv)/embed/[a-zA-Z0-9-_]+)',
+ webpage)
def _real_extract(self, url):
video_id = self._match_id(url)
if 'File not found' in webpage or 'deleted by the owner' in webpage:
raise ExtractorError('File not found', expected=True)
- # The following decryption algorithm is written by @yokrysty and
- # declared to be freely used in youtube-dl
- # See https://github.com/rg3/youtube-dl/issues/10408
- enc_data = self._html_search_regex(
- r'<span[^>]*>([^<]+)</span>\s*<span[^>]*>[^<]+</span>\s*<span[^>]+id="streamurl"',
- webpage, 'encrypted data')
-
- enc_code = self._html_search_regex(r'<script[^>]+>(゚ω゚[^<]+)</script>',
- webpage, 'encrypted code')
-
- js_code = self.openload_decode(enc_code)
- jsi = JSInterpreter(js_code)
-
- m_offset_fun = self._search_regex(r'slice\(0\s*-\s*(%s)\(\)' % _NAME_RE, js_code, 'javascript offset function')
- m_diff_fun = self._search_regex(r'charCodeAt\(0\)\s*\+\s*(%s)\(\)' % _NAME_RE, js_code, 'javascript diff function')
-
- offset = jsi.call_function(m_offset_fun)
- diff = jsi.call_function(m_diff_fun)
-
- video_url_chars = []
-
- for idx, c in enumerate(enc_data):
- j = compat_ord(c)
- if j >= 33 and j <= 126:
- j = ((j + 14) % 94) + 33
- if idx == len(enc_data) - offset:
- j += diff
- video_url_chars += compat_chr(j)
-
- video_url = 'https://openload.co/stream/%s?mime=true' % ''.join(video_url_chars)
+ ol_id = self._search_regex(
+ '<span[^>]+id="[^"]+"[^>]*>([0-9A-Za-z]+)</span>',
+ webpage, 'openload ID')
+
+ decoded = ''
+ a = ol_id[0:24]
+ b = []
+ for i in range(0, len(a), 8):
+ b.append(int(a[i:i + 8] or '0', 16))
+ ol_id = ol_id[24:]
+ j = 0
+ k = 0
+ while j < len(ol_id):
+ c = 128
+ d = 0
+ e = 0
+ f = 0
+ _more = True
+ while _more:
+ if j + 1 >= len(ol_id):
+ c = 143
+ f = int(ol_id[j:j + 2] or '0', 16)
+ j += 2
+ d += (f & 127) << e
+ e += 7
+ _more = f >= c
+ g = d ^ b[k % 3]
+ for i in range(4):
+ char_dec = (g >> 8 * i) & (c + 127)
+ char = compat_chr(char_dec)
+ if char != '#':
+ decoded += char
+ k += 1
+
+ video_url = 'https://openload.co/stream/%s?mime=true'
+ video_url = video_url % decoded
title = self._og_search_title(webpage, default=None) or self._search_regex(
r'<span[^>]+class=["\']title["\'][^>]*>([^<]+)', webpage,
'thumbnail': self._og_search_thumbnail(webpage, default=None),
'url': video_url,
# Seems all videos have extensions in their titles
- 'ext': determine_ext(title),
+ 'ext': determine_ext(title, 'mp4'),
'subtitles': subtitles,
}
-
return info_dict