diff --git a/src/extractor/Vidoza.ts b/src/extractor/Vidoza.ts new file mode 100644 index 0000000..7cc2a1a --- /dev/null +++ b/src/extractor/Vidoza.ts @@ -0,0 +1,103 @@ +import bytes from 'bytes'; +import { NotFoundError } from '../error'; +import { Context, Format, Meta, UrlResult } from '../types'; +import { + buildMediaFlowProxyExtractorStreamUrl, + supportsMediaFlowProxy, +} from '../utils'; +import { Extractor } from './Extractor'; + +/** @see https://github.com/Gujal00/ResolveURL/blob/master/script.module.resolveurl/lib/resolveurl/plugins/vidoza.py */ +export class Vidoza extends Extractor { + public readonly id = 'Vidoza'; + public readonly label = 'Vidoza (via MediaFlow Proxy)'; + public override readonly ttl = 10800000; // 3h + public override viaMediaFlowProxy = true; + + private domains = ['vidoza.net', 'vidoza.co', 'videzz.net']; + + public supports(ctx: Context, url: URL): boolean { + return ( + this.domains.some(d => url.host.includes(d)) + && supportsMediaFlowProxy(ctx) + ); + } + + public override normalize(url: URL): URL { + const id + = url.pathname.match(/embed-([A-Za-z0-9]+)\.html?/i)?.[1] + || url.pathname.match(/\/([A-Za-z0-9]+)\.html?/i)?.[1]; + + if (!id) return url; + + return new URL(`https://videzz.net/${id}.html`); + } + + protected override async extractInternal( + ctx: Context, + url: URL, + meta: Meta, + ): Promise { + const headers: Record = { + 'Referer': 'https://vidoza.net/', + 'User-Agent': + 'Mozilla/5.0 (Windows NT 10.0; Win64; x64) AppleWebKit/537.36 (KHTML, like Gecko) Chrome/120.0.0.0 Safari/537.36', + 'Accept': '*/*', + 'Accept-Language': 'en-US,en;q=0.9', + }; + + const html = await this.fetcher.text(ctx, url, headers); + if (!html) throw new NotFoundError('Vidoza: video page unavailable'); + + const titleMatch = html.match(/]*>([^<]+)<\/h1>/i); + const title = titleMatch?.[1]?.trim() || this.label; + + let bytesSize: number | null = null; + const sizeMatch = html.match(/File size:\s*([\d.]+\s*[GM]B)<\/span>/i); + + if (sizeMatch?.[1]) { + const parsed = bytes.parse(sizeMatch[1]); + if (parsed !== null) { + bytesSize = parsed; + } + } + + const extractedLabel = this.extractLabelFromHtml(html); + const height = extractedLabel ? parseInt(extractedLabel, 10) : null; + + const proxiedUrl = await buildMediaFlowProxyExtractorStreamUrl( + ctx, + this.fetcher, + this.id, + url, + headers, + ); + + return [ + { + url: proxiedUrl, + format: Format.mp4, + label: this.label, + ttl: this.ttl, + requestHeaders: headers, + sourceId: `${this.id}_${meta.countryCodes?.join('_') ?? 'all'}`, + meta: { + ...meta, + title, + ...(height !== null ? { height } : {}), + ...(bytesSize && bytesSize > 16777216 ? { bytes: bytesSize } : {}), + }, + }, + ]; + } + + private extractLabelFromHtml(html: string): string | null { + const regex + = /["']?\s*(?:file|src)\s*["']?\s*[:=,]?\s*["'][^"']+(?:[^}>\]]+)["']?\s*res\s*["']?\s*[:=]\s*["']?(\d{3,4})/i; + + const m + = html.match(regex); + + return m && m[1] ? m[1] : null; + } +}