fix: Quanben and Add Filters (#1877)
* fix quanben +added filters * remove https://69shu.biz
This commit is contained in:
@@ -1,201 +0,0 @@
|
||||
import { load as parseHTML } from 'cheerio';
|
||||
import { fetchText } from '@/lib/fetch';
|
||||
import { FilterTypes, Filters } from '@libs/filterInputs';
|
||||
import { Plugin } from '@/types/plugin';
|
||||
import { encode } from 'urlencode';
|
||||
import { NovelStatus } from '@libs/novelStatus';
|
||||
|
||||
class XinShu69 implements Plugin.PluginBase {
|
||||
id = '69xinshu';
|
||||
name = '69书吧';
|
||||
icon = 'src/cn/69xinshu/icon.png';
|
||||
site = 'https://69shu.biz';
|
||||
version = '0.1.2';
|
||||
|
||||
async popularNovels(
|
||||
pageNo: number,
|
||||
{ filters }: Plugin.PopularNovelsOptions<typeof this.filters>,
|
||||
): Promise<Plugin.NovelItem[]> {
|
||||
let url: string;
|
||||
if (pageNo === 1) {
|
||||
url = `${this.site}/novels/class/${filters.class.value}.htm`;
|
||||
} else {
|
||||
url = `${this.site}/ajax_novels/class/${filters.class.value}/${pageNo}.htm`;
|
||||
}
|
||||
|
||||
const body = await fetchText(url, {}, 'gbk');
|
||||
if (body === '') throw Error('无法获取小说列表,请检查网络');
|
||||
|
||||
const loadedCheerio = parseHTML(body);
|
||||
|
||||
const novels: Plugin.NovelItem[] = [];
|
||||
|
||||
const novelsList =
|
||||
pageNo === 1
|
||||
? loadedCheerio('.newlistbox > ul > li')
|
||||
: loadedCheerio('li');
|
||||
|
||||
novelsList.each((i, e) => {
|
||||
const novelUrl = loadedCheerio(e).find('div > h3 > a').attr('href');
|
||||
if (novelUrl) {
|
||||
const novelName = loadedCheerio(e).find('div > h3 > a').text();
|
||||
const novelCover = loadedCheerio(e).find('a > img').attr('data-src');
|
||||
|
||||
const novel = {
|
||||
name: novelName,
|
||||
cover: novelCover,
|
||||
path: novelUrl.replace(this.site, ''),
|
||||
};
|
||||
|
||||
novels.push(novel);
|
||||
}
|
||||
});
|
||||
|
||||
return novels;
|
||||
}
|
||||
|
||||
async parseNovel(novelPath: string): Promise<Plugin.SourceNovel> {
|
||||
const url = this.site + novelPath;
|
||||
|
||||
const body = await fetchText(url, {}, 'gbk');
|
||||
if (body === '') throw Error('无法获取小说内容,请检查网络');
|
||||
|
||||
const loadedCheerio = parseHTML(body);
|
||||
|
||||
const novel: Plugin.SourceNovel = {
|
||||
path: novelPath,
|
||||
chapters: [],
|
||||
name: loadedCheerio('h1 > a').text(),
|
||||
};
|
||||
|
||||
novel.cover = loadedCheerio('div.bookimg2 > img').attr('src');
|
||||
|
||||
novel.summary = loadedCheerio('div.navtxt').text().trim();
|
||||
|
||||
const bookInfo = loadedCheerio('div.booknav2').text();
|
||||
|
||||
const authorMatch = bookInfo.match(/作者:(.*)/);
|
||||
if (authorMatch) {
|
||||
novel.author = authorMatch[1];
|
||||
}
|
||||
|
||||
novel.artist = undefined;
|
||||
|
||||
novel.status = bookInfo.includes('连载')
|
||||
? NovelStatus.Ongoing
|
||||
: NovelStatus.Completed;
|
||||
|
||||
const genresMatch = bookInfo.match(/分类:(.*)/);
|
||||
if (genresMatch) {
|
||||
novel.genres = genresMatch[1];
|
||||
}
|
||||
|
||||
// Table of Content is on a different page than the summary page
|
||||
const chapters: Plugin.ChapterItem[] = [];
|
||||
|
||||
const chaptersUrl = loadedCheerio('a.more-btn').attr('href');
|
||||
if (chaptersUrl) {
|
||||
const chaptersBody = await fetchText(chaptersUrl, {}, 'gbk');
|
||||
|
||||
const chaptersLoadedCheerio = parseHTML(chaptersBody);
|
||||
|
||||
chaptersLoadedCheerio('li').each((i, e) => {
|
||||
const chapterUrl = chaptersLoadedCheerio(e).find('a').attr('href');
|
||||
|
||||
if (chapterUrl?.startsWith('https://')) {
|
||||
const chapterName = chaptersLoadedCheerio(e).find('a').text().trim();
|
||||
|
||||
chapters.push({
|
||||
name: chapterName,
|
||||
path: chapterUrl.replace(this.site, ''),
|
||||
});
|
||||
}
|
||||
});
|
||||
}
|
||||
|
||||
novel.chapters = chapters;
|
||||
|
||||
return novel;
|
||||
}
|
||||
|
||||
async parseChapter(chapterPath: string): Promise<string> {
|
||||
const body = await fetchText(this.site + chapterPath, {}, 'gbk');
|
||||
|
||||
const loadedCheerio = parseHTML(body);
|
||||
|
||||
const chapterText = (loadedCheerio('div.txtnav').prop('innerText') ?? '')
|
||||
.split('\n')
|
||||
// remove empty lines
|
||||
.map((line: string) => line.trim())
|
||||
.filter((line: string) => line !== '')
|
||||
// remove the first two lines which are the chapter name and author name
|
||||
.slice(2)
|
||||
.map((line: string) => `<p>${line}</p>`)
|
||||
.join('\n');
|
||||
|
||||
return chapterText;
|
||||
}
|
||||
|
||||
async searchNovels(
|
||||
searchTerm: string,
|
||||
pageNo: number,
|
||||
): Promise<Plugin.NovelItem[]> {
|
||||
if (pageNo > 1) return [];
|
||||
|
||||
const searchUrl = `${this.site}/modules/article/search.php`;
|
||||
const formData = new FormData();
|
||||
formData.append('searchkey', encode(searchTerm, 'gbk'));
|
||||
formData.append('searchtype', 'all');
|
||||
|
||||
const body = await fetchText(
|
||||
searchUrl,
|
||||
{ method: 'post', body: formData },
|
||||
'gbk',
|
||||
);
|
||||
if (body === '') throw Error('无法获取搜索结果,请检查网络');
|
||||
const loadedCheerio = parseHTML(body);
|
||||
|
||||
const novels: Plugin.NovelItem[] = [];
|
||||
|
||||
loadedCheerio('.newbox > ul > li').each((i, e) => {
|
||||
const novelUrl = loadedCheerio(e).find('a').attr('href')!;
|
||||
const novelName = loadedCheerio(e).find('div.newnav > h3').text();
|
||||
const novelCover = loadedCheerio(e).find('img').attr('data-src');
|
||||
|
||||
const novel = {
|
||||
name: novelName,
|
||||
path: novelUrl.replace(this.site, ''),
|
||||
cover: novelCover,
|
||||
};
|
||||
|
||||
novels.push(novel);
|
||||
});
|
||||
|
||||
return novels;
|
||||
}
|
||||
|
||||
filters = {
|
||||
class: {
|
||||
label: '分类',
|
||||
value: '0',
|
||||
options: [
|
||||
{ label: '全部分类', value: '0' },
|
||||
{ label: '言情小说', value: '3' },
|
||||
{ label: '玄幻魔法', value: '1' },
|
||||
{ label: '修真武侠', value: '2' },
|
||||
{ label: '穿越时空', value: '11' },
|
||||
{ label: '都市小说', value: '9' },
|
||||
{ label: '历史军事', value: '4' },
|
||||
{ label: '游戏竞技', value: '5' },
|
||||
{ label: '科幻空间', value: '6' },
|
||||
{ label: '悬疑惊悚', value: '7' },
|
||||
{ label: '同人小说', value: '8' },
|
||||
{ label: '官场职场', value: '10' },
|
||||
{ label: '青春校园', value: '12' },
|
||||
],
|
||||
type: FilterTypes.Picker,
|
||||
},
|
||||
} satisfies Filters;
|
||||
}
|
||||
|
||||
export default new XinShu69();
|
||||
+173
-309
@@ -1,53 +1,44 @@
|
||||
import { Plugin } from '@/types/plugin';
|
||||
import { FilterTypes, Filters } from '@libs/filterInputs';
|
||||
import { fetchApi, fetchFile } from '@/lib/fetch';
|
||||
import { fetchApi } from '@/lib/fetch';
|
||||
import { NovelStatus } from '@libs/novelStatus';
|
||||
import { load as parseHTML } from 'cheerio';
|
||||
import { defaultCover } from '@libs/defaultCover';
|
||||
|
||||
const getStandardNovelPath = (url: string | undefined): string | undefined => {
|
||||
const parseUrl = (url?: string): URL | undefined => {
|
||||
if (!url) return undefined;
|
||||
try {
|
||||
const parsedUrl = new URL(url);
|
||||
const match = parsedUrl.pathname.match(/^(\/amp)?(\/n\/[^\/]+\/)/);
|
||||
return match?.[2];
|
||||
} catch (error) {
|
||||
return new URL(url, 'https://www.quanben.io');
|
||||
} catch {
|
||||
return undefined;
|
||||
}
|
||||
};
|
||||
|
||||
const getChapterFileName = (url: string | undefined): string | undefined => {
|
||||
if (!url) return undefined;
|
||||
try {
|
||||
const parsedUrl = new URL(url);
|
||||
const pathParts = parsedUrl.pathname.split('/');
|
||||
const fileName = pathParts[pathParts.length - 1];
|
||||
if (fileName && /^\d+\.html$/.test(fileName)) {
|
||||
return fileName;
|
||||
}
|
||||
return undefined;
|
||||
} catch (error) {
|
||||
return undefined;
|
||||
}
|
||||
const getStandardNovelPath = (url?: string): string | undefined => {
|
||||
const parsedUrl = parseUrl(url);
|
||||
if (!parsedUrl) return undefined;
|
||||
const match = parsedUrl.pathname.match(/^(\/amp)?(\/n\/[^\/]+\/)/);
|
||||
return match?.[2];
|
||||
};
|
||||
|
||||
const getChapterFileName = (url?: string): string | undefined => {
|
||||
const parsedUrl = parseUrl(url);
|
||||
if (!parsedUrl) return undefined;
|
||||
const fileName = parsedUrl.pathname.split('/').pop();
|
||||
if (fileName && /^\d+\.html$/.test(fileName)) return fileName;
|
||||
return undefined;
|
||||
};
|
||||
|
||||
const makeAbsolute = (
|
||||
relativeUrl: string | undefined,
|
||||
baseUrl: string,
|
||||
relativeUrl?: string,
|
||||
baseUrl?: string,
|
||||
): string | undefined => {
|
||||
if (!relativeUrl) return undefined;
|
||||
if (!relativeUrl || !baseUrl) return undefined;
|
||||
try {
|
||||
if (relativeUrl.startsWith('//')) {
|
||||
return 'https:' + relativeUrl;
|
||||
}
|
||||
if (
|
||||
relativeUrl.startsWith('http://') ||
|
||||
relativeUrl.startsWith('https://')
|
||||
) {
|
||||
return relativeUrl;
|
||||
}
|
||||
if (relativeUrl.startsWith('//')) return 'https:' + relativeUrl;
|
||||
if (/^https?:\/\//.test(relativeUrl)) return relativeUrl;
|
||||
return new URL(relativeUrl, baseUrl).href;
|
||||
} catch (e) {
|
||||
} catch {
|
||||
return undefined;
|
||||
}
|
||||
};
|
||||
@@ -56,357 +47,230 @@ class QuanbenPlugin implements Plugin.PluginBase {
|
||||
id = 'quanben';
|
||||
name = 'Quanben';
|
||||
site = 'https://www.quanben.io/';
|
||||
version = '1.0.0';
|
||||
version = '1.0.1';
|
||||
icon = 'src/cn/quanben/icon.png';
|
||||
defaultCover = defaultCover;
|
||||
|
||||
filters = {} satisfies Filters;
|
||||
// filters
|
||||
filters = {
|
||||
genre: {
|
||||
label: '分类',
|
||||
value: 'all',
|
||||
options: [
|
||||
{ label: '全部', value: 'all' },
|
||||
{ label: '玄幻', value: 'xuanhuan' },
|
||||
{ label: '都市', value: 'dushi' },
|
||||
{ label: '言情', value: 'yanqing' },
|
||||
{ label: '穿越', value: 'chuanyue' },
|
||||
{ label: '青春', value: 'qingchun' },
|
||||
{ label: '仙侠', value: 'xianxia' },
|
||||
{ label: '灵异', value: 'lingyi' },
|
||||
{ label: '悬疑', value: 'xuanyi' },
|
||||
{ label: '历史', value: 'lishi' },
|
||||
{ label: '军事', value: 'junshi' },
|
||||
{ label: '游戏', value: 'youxi' },
|
||||
{ label: '竞技', value: 'jingji' },
|
||||
{ label: '科幻', value: 'kehuan' },
|
||||
{ label: '职场', value: 'zhichang' },
|
||||
{ label: '官场', value: 'guanchang' },
|
||||
{ label: '现言', value: 'xianyan' },
|
||||
{ label: '耽美', value: 'danmei' },
|
||||
{ label: '其它', value: 'qita' },
|
||||
],
|
||||
type: FilterTypes.Picker,
|
||||
},
|
||||
} satisfies Filters;
|
||||
|
||||
async popularNovels(_pageNo: number): Promise<Plugin.NovelItem[]> {
|
||||
const url = this.site + 'amp/';
|
||||
const result = await fetchApi(url);
|
||||
if (!result.ok) {
|
||||
throw new Error(
|
||||
`[Quanben] Failed to fetch AMP popular novels page: ${url} - Status: ${result.status}`,
|
||||
);
|
||||
}
|
||||
const body = await result.text();
|
||||
const $ = parseHTML(body);
|
||||
// homepage, when you first open the extension (with the applied filters if any)
|
||||
async popularNovels(
|
||||
_pageNo: number,
|
||||
{ filters }: Plugin.PopularNovelsOptions<typeof this.filters>,
|
||||
): Promise<Plugin.NovelItem[]> {
|
||||
const url =
|
||||
filters.genre.value === 'all'
|
||||
? this.site
|
||||
: `${this.site}c/${filters.genre.value}.html`;
|
||||
|
||||
const res = await fetchApi(url);
|
||||
if (!res.ok)
|
||||
throw new Error(`[Quanben] Failed to fetch: ${url} - ${res.status}`);
|
||||
|
||||
const $ = parseHTML(await res.text());
|
||||
const novels: Plugin.NovelItem[] = [];
|
||||
const processedAmpPaths = new Set<string>();
|
||||
|
||||
$('div.box').each((_i, box) => {
|
||||
const $box = $(box);
|
||||
$('div.list2').each((_i, list2) => {
|
||||
const $list2 = $(list2);
|
||||
const $link = $list2.find('h3 > a').first();
|
||||
const href = $link.attr('href')?.trim();
|
||||
const name = $link.text().trim();
|
||||
const rawCover =
|
||||
$list2.find('img').attr('src')?.trim() ||
|
||||
$list2.find('img').attr('data-src')?.trim();
|
||||
const cover = makeAbsolute(rawCover, this.site) || this.defaultCover;
|
||||
|
||||
// 1. Process the featured novel (div.list2) if it exists
|
||||
const $featured = $box.find('div.list2');
|
||||
if ($featured.length > 0) {
|
||||
const $link = $featured.find('h3 > a');
|
||||
const ampPath = $link.attr('href')?.trim();
|
||||
const name = $link.text().trim();
|
||||
const rawCoverSrc = $featured.find('amp-img').attr('src')?.trim();
|
||||
const cover = makeAbsolute(rawCoverSrc, this.site) || this.defaultCover;
|
||||
|
||||
if (ampPath && name && !processedAmpPaths.has(ampPath)) {
|
||||
novels.push({ name, path: ampPath, cover });
|
||||
processedAmpPaths.add(ampPath);
|
||||
}
|
||||
if (href && name) {
|
||||
const path = getStandardNovelPath(href);
|
||||
if (path) novels.push({ name, path, cover });
|
||||
}
|
||||
});
|
||||
|
||||
// 2. Process novels in the list (ul.list)
|
||||
$box.find('ul.list li').each((_j, listItem) => {
|
||||
const $listItem = $(listItem);
|
||||
const $link = $listItem.find('a');
|
||||
const ampPath = $link.attr('href')?.trim();
|
||||
const name = $link.find('span').first().text().trim(); // Name is inside a span within the link
|
||||
// only first entry bcs the others dont have an image
|
||||
$('ul.list').each((_i, ul) => {
|
||||
const $firstLi = $(ul).find('li').first();
|
||||
const $a = $firstLi.find('a').first();
|
||||
const href = $a.attr('href')?.trim();
|
||||
const name =
|
||||
$a.text().trim() || $firstLi.find('span.author').text().trim();
|
||||
const cover = this.defaultCover;
|
||||
|
||||
if (ampPath && name && !processedAmpPaths.has(ampPath)) {
|
||||
novels.push({ name, path: ampPath, cover: this.defaultCover }); // Use default cover for list items
|
||||
processedAmpPaths.add(ampPath);
|
||||
}
|
||||
});
|
||||
if (href && name) {
|
||||
const path = getStandardNovelPath(href);
|
||||
if (path) novels.push({ name, path, cover });
|
||||
}
|
||||
});
|
||||
|
||||
return novels;
|
||||
}
|
||||
|
||||
// novel details and metadata
|
||||
async parseNovel(novelPath: string): Promise<Plugin.SourceNovel> {
|
||||
// ** Expects AMP path: /amp/n/novel-name/ **
|
||||
if (
|
||||
!novelPath ||
|
||||
!novelPath.startsWith('/amp/n/') ||
|
||||
!novelPath.endsWith('/')
|
||||
) {
|
||||
throw new Error(
|
||||
`[Quanben parseNovel] Invalid novelPath received: "${novelPath}". Expected AMP format: "/amp/n/novel-name/"`,
|
||||
);
|
||||
}
|
||||
const standardPath = novelPath.replace(/^\/amp/, '');
|
||||
if (!standardPath.startsWith('/n/') || !standardPath.endsWith('/'))
|
||||
throw new Error(`[Quanben parseNovel] Invalid path: ${novelPath}`);
|
||||
|
||||
// Construct full AMP novel URL
|
||||
const fullNovelUrl = makeAbsolute(novelPath, this.site);
|
||||
if (!fullNovelUrl) {
|
||||
throw new Error(
|
||||
`[Quanben parseNovel] Could not construct full AMP novel URL from path: ${novelPath}`,
|
||||
);
|
||||
}
|
||||
const fullUrl = makeAbsolute(standardPath, this.site);
|
||||
if (!fullUrl)
|
||||
throw new Error(`[Quanben parseNovel] Could not construct full URL`);
|
||||
|
||||
// 1. Fetch AMP Novel Page HTML
|
||||
const novelPageResult = await fetchApi(fullNovelUrl);
|
||||
if (!novelPageResult.ok) {
|
||||
throw new Error(
|
||||
`[Quanben parseNovel] Failed to fetch AMP novel page: ${fullNovelUrl} - Status: ${novelPageResult.status}`,
|
||||
);
|
||||
}
|
||||
const novelPageHtml = await novelPageResult.text();
|
||||
const $ = parseHTML(novelPageHtml);
|
||||
const res = await fetchApi(fullUrl);
|
||||
if (!res.ok)
|
||||
throw new Error(`[Quanben parseNovel] Failed to fetch: ${fullUrl}`);
|
||||
|
||||
// 2. Parse Novel Details from AMP page - **Refined Selectors**
|
||||
const $infoBox = $('div.list2'); // Base element for many details
|
||||
const $descriptionBox = $('div.description'); // Separate element for summary
|
||||
const $ = parseHTML(await res.text());
|
||||
const $info = $('div.list2').first();
|
||||
const $desc = $('div.description').first();
|
||||
|
||||
const novel: Plugin.SourceNovel = {
|
||||
path: novelPath, // Use the input AMP path
|
||||
name:
|
||||
$infoBox.find('h3').text().trim() || // **Use h3 inside list2**
|
||||
$('h1[itemprop="name headline"]').text().trim() || // Fallback H1
|
||||
'Unknown Novel Name',
|
||||
path: standardPath,
|
||||
name: $info.find('h3').text().trim() || 'Unknown Novel',
|
||||
cover:
|
||||
makeAbsolute($infoBox.find('amp-img').attr('src'), this.site) || // **amp-img inside list2**
|
||||
makeAbsolute($info.find('img').attr('src'), this.site) ||
|
||||
this.defaultCover,
|
||||
summary:
|
||||
$descriptionBox.find('p').text().trim() || // **Use p inside div.description**
|
||||
$descriptionBox.text().trim() || // Fallback to div.description text
|
||||
undefined,
|
||||
// Use :contains() for more robust selection within the info box
|
||||
author:
|
||||
$infoBox.find("p:contains('作者:') span").text().trim() || undefined,
|
||||
status: NovelStatus.Unknown, // Parsed below
|
||||
genres:
|
||||
$infoBox.find("p:contains('类别:') span").text().trim() || undefined,
|
||||
chapters: [], // Parsed below
|
||||
$desc.find('p').text().trim() || $desc.text().trim() || undefined,
|
||||
author: $info.find("p:contains('作者:') span").text().trim() || undefined,
|
||||
status: NovelStatus.Unknown,
|
||||
genres: $info.find("p:contains('类别:') span").text().trim() || undefined,
|
||||
chapters: await this.parseChapterList(standardPath),
|
||||
};
|
||||
|
||||
// Parse Status from AMP page
|
||||
const statusText =
|
||||
$infoBox.find("p:contains('状态:') span").text().trim() || // Try specific span first
|
||||
$infoBox.text(); // Fallback to full text
|
||||
if (statusText.includes('完结') || statusText.includes('已完成')) {
|
||||
novel.status = NovelStatus.Completed;
|
||||
} else if (statusText.includes('连载中') || statusText.includes('进行中')) {
|
||||
novel.status = NovelStatus.Ongoing;
|
||||
}
|
||||
|
||||
// 3. Fetch and Parse Chapter List
|
||||
novel.chapters = await this.parseChapterList(novelPath);
|
||||
novel.status = NovelStatus.Completed;
|
||||
|
||||
return novel;
|
||||
}
|
||||
|
||||
// Separate function to handle chapter list parsing - **Using AMP**
|
||||
async parseChapterList(novelPath: string): Promise<Plugin.ChapterItem[]> {
|
||||
// ** Expects AMP path: /amp/n/novel-name/ **
|
||||
if (
|
||||
!novelPath ||
|
||||
!novelPath.startsWith('/amp/n/') ||
|
||||
!novelPath.endsWith('/')
|
||||
) {
|
||||
return [];
|
||||
}
|
||||
if (!novelPath.startsWith('/n/') || !novelPath.endsWith('/')) return [];
|
||||
|
||||
// Construct the AMP chapter list URL
|
||||
const ampChapterListUrl = makeAbsolute(novelPath + 'list.html', this.site);
|
||||
if (!ampChapterListUrl) {
|
||||
return [];
|
||||
}
|
||||
const url = makeAbsolute(novelPath + 'list.html', this.site);
|
||||
if (!url) return [];
|
||||
|
||||
const chapterListResult = await fetchApi(ampChapterListUrl);
|
||||
if (!chapterListResult.ok) {
|
||||
return [];
|
||||
}
|
||||
const chapterListHtml = await chapterListResult.text();
|
||||
const $ = parseHTML(chapterListHtml);
|
||||
const res = await fetchApi(url);
|
||||
if (!res.ok) return [];
|
||||
|
||||
const $ = parseHTML(await res.text());
|
||||
const chapters: Plugin.ChapterItem[] = [];
|
||||
// Extract standard novel name from AMP path for chapter path storage
|
||||
const standardNovelPathMatch = novelPath.match(/(\/n\/[^\/]+\/)/);
|
||||
if (!standardNovelPathMatch || !standardNovelPathMatch[1]) {
|
||||
return [];
|
||||
}
|
||||
const novelNameOnly = standardNovelPathMatch[1].replace(/^\/n\/|\/$/g, '');
|
||||
const novelName = novelPath.match(/\/n\/([^\/]+)\//)?.[1];
|
||||
if (!novelName) return [];
|
||||
|
||||
// 1. Parse chapters from ALL list3 uls in the AMP HTML
|
||||
$('ul.list3 li a').each((_i, el) => {
|
||||
const $el = $(el);
|
||||
const chapterName = $el.text().trim();
|
||||
const chapterHref = $el.attr('href');
|
||||
if (chapterName && chapterHref) {
|
||||
// URLs in AMP page usually point to the standard chapter URLs
|
||||
const absoluteUrl = makeAbsolute(chapterHref, this.site); // Use base site URL
|
||||
const chapterFileName = getChapterFileName(absoluteUrl);
|
||||
if (chapterFileName) {
|
||||
// ** Store chapter path in standard format for parseChapter **
|
||||
const chapterPathForStorage = novelNameOnly + '/' + chapterFileName;
|
||||
chapters.push({
|
||||
name: chapterName,
|
||||
path: chapterPathForStorage,
|
||||
// chapterNumber assigned later
|
||||
});
|
||||
}
|
||||
}
|
||||
const name = $el.text().trim();
|
||||
const href = $el.attr('href');
|
||||
if (!name || !href) return;
|
||||
|
||||
const fileName = getChapterFileName(makeAbsolute(href, this.site));
|
||||
if (!fileName) return;
|
||||
|
||||
chapters.push({
|
||||
name,
|
||||
path: `${novelName}/${fileName}`,
|
||||
});
|
||||
});
|
||||
|
||||
// 2. Deduplicate chapters based on path
|
||||
const chapterMap = new Map<string, Plugin.ChapterItem>();
|
||||
chapters.forEach(chapter => {
|
||||
if (!chapterMap.has(chapter.path)) {
|
||||
chapterMap.set(chapter.path, chapter);
|
||||
}
|
||||
});
|
||||
const uniqueChapters = Array.from(chapterMap.values());
|
||||
|
||||
// 3. Sort chapters numerically based on filename
|
||||
const uniqueChapters = Array.from(
|
||||
new Map(chapters.map(c => [c.path, c])).values(),
|
||||
);
|
||||
uniqueChapters.sort((a, b) => {
|
||||
const numA = parseInt(a.path.match(/(\d+)\.html$/)?.[1] || '0', 10);
|
||||
const numB = parseInt(b.path.match(/(\d+)\.html$/)?.[1] || '0', 10);
|
||||
return numA - numB;
|
||||
});
|
||||
|
||||
// 4. Assign chapter numbers
|
||||
return uniqueChapters.map((chapter, index) => ({
|
||||
...chapter,
|
||||
chapterNumber: index + 1,
|
||||
}));
|
||||
return uniqueChapters.map((c, idx) => ({ ...c, chapterNumber: idx + 1 }));
|
||||
}
|
||||
|
||||
async parseChapter(chapterPath: string): Promise<string> {
|
||||
// ** Expects standard path: novel-name/chapterFileName.html **
|
||||
if (
|
||||
!chapterPath ||
|
||||
!chapterPath.includes('/') ||
|
||||
chapterPath.endsWith('/')
|
||||
) {
|
||||
throw new Error(
|
||||
`[Quanben] Invalid chapterPath format received in parseChapter: "${chapterPath}". Expected format: "novel-name/chapterFileName.html"`,
|
||||
);
|
||||
}
|
||||
if (!chapterPath.includes('/') || chapterPath.endsWith('/'))
|
||||
throw new Error(`[Quanben] Invalid chapter path: "${chapterPath}"`);
|
||||
|
||||
// Construct the standard chapter URL using /n/ prefix
|
||||
const chapterUrl = `${this.site}n/${chapterPath}`;
|
||||
const url = `${this.site}n/${chapterPath}`;
|
||||
const res = await fetchApi(url);
|
||||
if (!res.ok) throw new Error(`[Quanben] Failed to fetch chapter: ${url}`);
|
||||
|
||||
const result = await fetchApi(chapterUrl);
|
||||
if (!result.ok) {
|
||||
// Handle potential redirects (e.g., chapter moved) - check Location header if status is 3xx
|
||||
if (result.status >= 300 && result.status < 400) {
|
||||
const redirectUrl = result.headers.get('Location');
|
||||
if (redirectUrl) {
|
||||
const absoluteRedirectUrl = makeAbsolute(redirectUrl, chapterUrl);
|
||||
if (!absoluteRedirectUrl) {
|
||||
throw new Error(
|
||||
`[Quanben] Failed to make redirected URL absolute: ${redirectUrl}`,
|
||||
);
|
||||
}
|
||||
const redirectResult = await fetchApi(absoluteRedirectUrl);
|
||||
if (!redirectResult.ok) {
|
||||
throw new Error(
|
||||
`[Quanben] Failed to fetch redirected chapter content: ${absoluteRedirectUrl} - Status: ${redirectResult.status}`,
|
||||
);
|
||||
}
|
||||
return this.extractChapterContent(
|
||||
await redirectResult.text(),
|
||||
absoluteRedirectUrl,
|
||||
);
|
||||
}
|
||||
}
|
||||
throw new Error(
|
||||
`[Quanben] Failed to fetch chapter content: ${chapterUrl} - Status: ${result.status}`,
|
||||
);
|
||||
}
|
||||
const body = await result.text();
|
||||
return this.extractChapterContent(body, chapterUrl);
|
||||
return this.extractChapterContent(await res.text());
|
||||
}
|
||||
|
||||
// Helper function to extract and clean chapter content from HTML body
|
||||
private extractChapterContent(body: string, urlForLog: string): string {
|
||||
private extractChapterContent(body: string): string {
|
||||
const $ = parseHTML(body);
|
||||
let $content = $('#contentbody');
|
||||
if (!$content.length) {
|
||||
$content = $('#content');
|
||||
}
|
||||
if (!$content.length) {
|
||||
$content = $('.content');
|
||||
}
|
||||
if (!$content.length) {
|
||||
return 'Error: Could not find chapter content container.';
|
||||
}
|
||||
let $content = $('#contentbody, #content, .content').first();
|
||||
if (!$content.length) return 'Error: Chapter content not found.';
|
||||
|
||||
$content
|
||||
.find(
|
||||
'script, style, ins, iframe, [class*="ads"], [id*="ads"], [class*="google"], [id*="google"], [class*="recommend"], div[align="center"]',
|
||||
)
|
||||
.remove();
|
||||
$content.find('p').each((_i, el) => {
|
||||
const $p = $(el);
|
||||
const pText = $p.text().trim();
|
||||
if (
|
||||
pText.includes('请记住本书首发域名') ||
|
||||
pText.includes('手机版阅读网址') ||
|
||||
pText.includes('quanben') ||
|
||||
pText.includes('最新网址') ||
|
||||
pText.includes('章节报错') ||
|
||||
pText.match(/app|APP|下载|客户端/) ||
|
||||
pText.length === 0 ||
|
||||
($p
|
||||
.html()
|
||||
?.replace(/ /g, '')
|
||||
.trim() === '' &&
|
||||
$p.find('img').length === 0)
|
||||
) {
|
||||
$p.remove();
|
||||
}
|
||||
});
|
||||
$content
|
||||
.contents()
|
||||
.filter(function () {
|
||||
return this.type === 'comment';
|
||||
})
|
||||
.remove();
|
||||
let chapterText = $content.html();
|
||||
if (!chapterText) {
|
||||
return 'Error: Chapter content was empty after cleaning.';
|
||||
}
|
||||
chapterText = chapterText.replace(/<\s*p[^>]*>/gi, '\n\n');
|
||||
chapterText = chapterText.replace(/<\s*br[^>]*>/gi, '\n');
|
||||
chapterText = chapterText.replace(/<[^>]+>/g, '');
|
||||
chapterText = parseHTML(`<div>${chapterText}</div>`).text();
|
||||
chapterText = chapterText.replace(/[\t ]+/g, ' ');
|
||||
chapterText = chapterText.replace(/\n{3,}/g, '\n\n');
|
||||
chapterText = chapterText.trim();
|
||||
return chapterText;
|
||||
|
||||
return (
|
||||
($content.html() || '').replace(/[\t ]+/g, ' ').trim() ||
|
||||
'Error: Chapter content empty.'
|
||||
);
|
||||
}
|
||||
|
||||
async searchNovels(
|
||||
searchTerm: string,
|
||||
_pageNo: number,
|
||||
): Promise<Plugin.NovelItem[]> {
|
||||
const searchUrl = `${this.site}index.php?c=book&a=search&keywords=${encodeURIComponent(searchTerm)}`;
|
||||
const result = await fetchApi(searchUrl);
|
||||
if (!result.ok) {
|
||||
return [];
|
||||
}
|
||||
const body = await result.text();
|
||||
const $ = parseHTML(body);
|
||||
// add search
|
||||
async searchNovels(searchTerm: string): Promise<Plugin.NovelItem[]> {
|
||||
const url = `${this.site}index.php?c=book&a=search&keywords=${encodeURIComponent(searchTerm)}`;
|
||||
const res = await fetchApi(url);
|
||||
if (!res.ok) return [];
|
||||
|
||||
const $ = parseHTML(await res.text());
|
||||
const novels: Plugin.NovelItem[] = [];
|
||||
|
||||
$('div.list2').each((_i, element) => {
|
||||
const $el = $(element);
|
||||
const nameLink = $el.find('h3 > a').first();
|
||||
const img = $el.find('img').first();
|
||||
$('div.list2').each((_i, el) => {
|
||||
const $el = $(el);
|
||||
const $link = $el.find('h3 > a').first();
|
||||
const href = $link.attr('href');
|
||||
const name = $link.text().trim();
|
||||
const cover = makeAbsolute(
|
||||
$el.find('img').attr('src') || $el.find('img').attr('data-src'),
|
||||
this.site,
|
||||
);
|
||||
|
||||
const novelName = nameLink.text().trim();
|
||||
const novelHref = nameLink.attr('href');
|
||||
const novelCover = img.attr('src') || img.attr('data-src');
|
||||
|
||||
if (novelHref && novelName) {
|
||||
const absoluteUrl = makeAbsolute(novelHref, this.site);
|
||||
const standardPath = getStandardNovelPath(absoluteUrl);
|
||||
|
||||
if (standardPath) {
|
||||
// **Construct AMP Path for storage**
|
||||
const ampPath = '/amp' + standardPath;
|
||||
const absoluteCover = makeAbsolute(novelCover, this.site);
|
||||
if (href && name) {
|
||||
const path = getStandardNovelPath(makeAbsolute(href, this.site));
|
||||
if (path)
|
||||
novels.push({
|
||||
name: novelName,
|
||||
path: ampPath, // **Store the AMP path**
|
||||
cover: absoluteCover || this.defaultCover,
|
||||
name,
|
||||
path: '/amp' + path,
|
||||
cover: cover || this.defaultCover,
|
||||
});
|
||||
}
|
||||
}
|
||||
});
|
||||
return novels;
|
||||
}
|
||||
|
||||
// Use fetchApi for fetchImage as it handles potential errors and returns Response
|
||||
async fetchImage(url: string): Promise<Response> {
|
||||
return fetchApi(url);
|
||||
}
|
||||
|
||||
Reference in New Issue
Block a user