X-Git-Url: http://git.bitcoin.ninja/index.cgi?a=blobdiff_plain;f=youtube_dl%2Fextractor%2Fwimp.py;h=3dab9145ba9c57bfd1d78a90a847761c23f0d8a8;hb=c11485162bcbf6f517fd9225850a048fc5f7cec1;hp=9d52c947ea4c96a011f612877b5bdc4e9c7b1826;hpb=d1bd37deac4fb61b350ba488bdbdf230aa0b335c;p=youtube-dl diff --git a/youtube_dl/extractor/wimp.py b/youtube_dl/extractor/wimp.py index 9d52c947e..3dab9145b 100644 --- a/youtube_dl/extractor/wimp.py +++ b/youtube_dl/extractor/wimp.py @@ -1,25 +1,58 @@ -import re -import base64 +from __future__ import unicode_literals + from .common import InfoExtractor +from .youtube import YoutubeIE class WimpIE(InfoExtractor): - _VALID_URL = r'(?:http://)?(?:www\.)?wimp\.com/([^/]+)/' + _VALID_URL = r'https?://(?:www\.)?wimp\.com/(?P[^/]+)' + _TESTS = [{ + 'url': 'http://www.wimp.com/maru-is-exhausted/', + 'md5': 'ee21217ffd66d058e8b16be340b74883', + 'info_dict': { + 'id': 'maru-is-exhausted', + 'ext': 'mp4', + 'title': 'Maru is exhausted.', + 'description': 'md5:57e099e857c0a4ea312542b684a869b8', + } + }, { + 'url': 'http://www.wimp.com/clowncar/', + 'md5': '5c31ad862a90dc5b1f023956faec13fe', + 'info_dict': { + 'id': 'cG4CEr2aiSg', + 'ext': 'webm', + 'title': 'Basset hound clown car...incredible!', + 'description': '5 of my Bassets crawled in this dog loo! www.bellinghambassets.com\n\nFor licensing/usage please contact: licensing(at)jukinmediadotcom', + 'upload_date': '20140303', + 'uploader': 'Gretchen Hoey', + 'uploader_id': 'gretchenandjeff1', + }, + 'add_ie': ['Youtube'], + }] def _real_extract(self, url): - mobj = re.match(self._VALID_URL, url) - video_id = mobj.group(1) + video_id = self._match_id(url) + webpage = self._download_webpage(url, video_id) - title = re.search('\',webpage).group(1) - thumbnail_url = re.search('\',webpage).group(1) - googleString = re.search("googleCode = '(.*?)'", webpage) - googleString = base64.b64decode(googleString.group(1)) - final_url = re.search('","(.*?)"', googleString).group(1) - ext = final_url.split('.')[-1] - return [{ - 'id': video_id, - 'url': final_url, - 'ext': ext, - 'title': title, - 'thumbnail': thumbnail_url, - }] + + youtube_id = self._search_regex( + (r"videoId\s*:\s*[\"']([0-9A-Za-z_-]{11})[\"']", + r'data-id=["\']([0-9A-Za-z_-]{11})'), + webpage, 'video URL', default=None) + if youtube_id: + return { + '_type': 'url', + 'url': youtube_id, + 'ie_key': YoutubeIE.ie_key(), + } + + info_dict = self._extract_jwplayer_data( + webpage, video_id, require_title=False) + + info_dict.update({ + 'id': video_id, + 'title': self._og_search_title(webpage), + 'description': self._og_search_description(webpage), + }) + + return info_dict