打开/关闭菜单
60
73
29
1.1K
武外梗百科
打开/关闭外观设置菜单
打开/关闭个人菜单
未登录
未登录用户的IP地址会在进行任意编辑后公开展示。

Module:PageSummary:修订间差异

武外梗百科 爱国好学自强图新的百科全书
无编辑摘要
无编辑摘要
第1行: 第1行:
local p = {}
local p = {}


-- stripImages 函数保持不变(已使用字节级,但对图片提取影响小)
-- ────────────────────────────────────────────────
-- 剥离图片、分类等 [[File:…]] [[Category:…]] 结构,避免出现在摘要中
-- ────────────────────────────────────────────────
local function stripImages(text)
    if not text or text == "" then return "" end
   
    -- 统一常见命名空间写法,减少匹配分支
    text = text:gsub("%[%[(文件|File|Image|图像):", "[[File:")
          :gsub("%[%[(分类|Category):", "[[Category:")
   
    local out = {}
    local i = 1
    local len = #text
   
    while i <= len do
        if text:byte(i) == 91 and text:byte(i+1) == 91 then  -- [[
            local chunk = text:sub(i, i+15)
            if chunk:match("^%[%[File:") or chunk:match("^%[%[Category:") then
                -- 跳过整段嵌套链接 [[ ... ]]
                local depth = 1
                local j = i + 2
                while j <= len and depth > 0 do
                    if text:byte(j) == 91 and text:byte(j+1) == 91 then
                        depth = depth + 1
                        j = j + 2
                    elseif text:byte(j) == 93 and text:byte(j+1) == 93 then
                        depth = depth - 1
                        j = j + 2
                    else
                        j = j + 1
                    end
                end
                i = j  -- 跳到 ]] 之后或文本末尾
            else
                -- 普通内部链接,保留
                table.insert(out, "[[")
                i = i + 2
            end
        else
            table.insert(out, text:sub(i,i))
            i = i + 1
        end
       
        if i > 1200 then break end  -- 极端安全阀
    end
   
    return table.concat(out)
end


-- cleanFinal 函数保持不变
-- ────────────────────────────────────────────────
-- 清理模板、参考文献、标题、列表等,得到纯文本摘要
-- ────────────────────────────────────────────────
local function cleanFinal(text)
    if not text or text == "" then return "" end
   
    text = text
        :gsub("{{[^{}]*}}", "")                  -- 模板(简单版,够用)
        :gsub("<ref.-</?ref[^>]*>", "")          -- 参考文献
        :gsub("<!%-%-.-%-%->", "")              -- HTML 注释
        :gsub("\n==+[^=]+==+", " ")              -- 各级标题
        :gsub("\n%s*[*#:;]+", " ")              -- 列表、定义列表
        :gsub("%s*\n%s*", " ")                  -- 所有换行 → 单个空格
        :gsub("%s+", " ")                        -- 压缩连续空格
   
    return mw.text.trim(text)
end


-- ────────────────────────────────────────────────
-- 主函数:生成带首图的简短摘要 + 标题
-- ────────────────────────────────────────────────
function p.getSummaryAndImage(frame)
function p.getSummaryAndImage(frame)
     local pageName = mw.text.trim(frame.args[1] or "")
     local pageName = mw.text.trim(frame.args[1] or "")
第15行: 第81行:
     if not content or content == "" then return "" end
     if not content or content == "" then return "" end
      
      
     -- 用 mw.ustring 安全截取(避免字节切割中文)
     -- 安全取前 300 个 unicode 字符(避免字节截断中文)
     local rawExcerpt = mw.ustring.sub(content, 1, 300) -- 字符级 300 个字符,够用
     local rawExcerpt = mw.ustring.sub(content, 1, 300)
      
      
     -- 改用 mw.ustring 的模式匹配,更安全处理中文文件名
     -- 提取第一张图片(优先标准写法)
     local firstImage
     local firstImage = mw.ustring.match(rawExcerpt, "%[%[(File|文件|Image|图像):([^|%]]+)")
    -- 尝试匹配 File: 或 文件: 开头的
    firstImage = mw.ustring.match(rawExcerpt, "%[%[(File|文件|Image|图像):([^|%]]+)")
      
      
     -- 如果没匹配到,再试更宽松的(但优先上面)
     -- 备选:更宽松匹配,但只接受文件类
     if not firstImage then
     if not firstImage then
         firstImage = mw.ustring.match(rawExcerpt, "%[%[([^|%]]+)%]")
         local candidate = mw.ustring.match(rawExcerpt, "%[%[([^|%]]+)%]")
         if firstImage and not (firstImage:find("File:") or firstImage:find("文件:")
         if candidate then
            or firstImage:find("Image:") or firstImage:find("图像:")) then
            local ns = candidate:match("^([^:]+):")
            firstImage = nil  -- 不是文件就丢掉
            if ns and (ns:lower() == "file" or ns == "文件" or ns:lower() == "image" or ns == "图像") then
                firstImage = candidate
            end
         end
         end
     end
     end
第36行: 第102行:
      
      
     local targetLen = 120
     local targetLen = 120
     local summary = cleanText:sub(1, targetLen * 2) or ""
     local summary   = cleanText:sub(1, targetLen * 2) or ""
      
      
     local ulen = mw.ustring.len(summary) or 0
     local ulen = mw.ustring.len(summary) or 0
第42行: 第108行:
         summary = mw.ustring.sub(summary, 1, targetLen) or ""
         summary = mw.ustring.sub(summary, 1, targetLen) or ""
          
          
        -- 尽量避免截在中文单词/句子中间
         local rev = summary:reverse()
         local rev = summary:reverse()
         local lastBreak = rev:find("[ %s%p,。!?;:]") or 1
         local lastBreak = rev:find("[ %s%p,。!?;:,…]") or 1
         if lastBreak > 1 and lastBreak < 20 then
         if lastBreak > 1 and lastBreak < 20 then
             summary = mw.ustring.sub(summary, 1, #summary - lastBreak + 1) or ""
             summary = mw.ustring.sub(summary, 1, #summary - lastBreak + 1) or ""
第53行: 第120行:
     summary = mw.text.trim(summary)
     summary = mw.text.trim(summary)
      
      
     -- HTML 构建
     -- ── HTML 输出 ───────────────────────────────────
     local container = mw.html.create('div')
     local container = mw.html.create('div')
         :css('margin-bottom', '25px')
         :css('margin-bottom', '25px')
         :css('display', 'flow-root')
         :css('display', 'flow-root')
      
      
    -- 标题
     container:tag('div')
     container:tag('div')
         :css('font-size', '1.3em')
         :css('font-size', '1.3em')
第65行: 第133行:
         :wikitext('[[' .. pageName .. ']]')
         :wikitext('[[' .. pageName .. ']]')
      
      
    -- 内容区(右浮图 + 文字)
     local textDiv = container:tag('div')
     local textDiv = container:tag('div')
         :css('line-height', '1.6')
         :css('line-height', '1.6')
第70行: 第139行:
         :css('font-size', '14px')
         :css('font-size', '14px')
      
      
    -- 关键修复:如果有图片名,确保用 mw.ustring 处理
     if firstImage and firstImage ~= "" then
     if firstImage and firstImage ~= "" then
        -- 额外清理:去掉可能的前后空格(虽然 match 已捕获干净)
         firstImage = mw.text.trim(firstImage)
         firstImage = mw.text.trim(firstImage)
       
         local img = '[[' .. firstImage .. '|120px|right|link=' .. pageName .. ']]'
        -- 用 [[File:xxx|...]] 格式,确保 MediaWiki 正确解析
         textDiv:wikitext(img)
         local imgWikitext = '[[' .. firstImage .. '|120px|right|link=' .. pageName .. ']]'
         textDiv:wikitext(imgWikitext)
     end
     end
      
      

2026年2月17日 (二) 23:54的版本

此模块的文档可以在Module:PageSummary/doc创建

local p = {}

-- ────────────────────────────────────────────────
-- 剥离图片、分类等 [[File:…]] [[Category:…]] 结构,避免出现在摘要中
-- ────────────────────────────────────────────────
local function stripImages(text)
    if not text or text == "" then return "" end
    
    -- 统一常见命名空间写法,减少匹配分支
    text = text:gsub("%[%[(文件|File|Image|图像):", "[[File:")
           :gsub("%[%[(分类|Category):", "[[Category:")
    
    local out = {}
    local i = 1
    local len = #text
    
    while i <= len do
        if text:byte(i) == 91 and text:byte(i+1) == 91 then  -- [[
            local chunk = text:sub(i, i+15)
            if chunk:match("^%[%[File:") or chunk:match("^%[%[Category:") then
                -- 跳过整段嵌套链接 [[ ... ]]
                local depth = 1
                local j = i + 2
                while j <= len and depth > 0 do
                    if text:byte(j) == 91 and text:byte(j+1) == 91 then
                        depth = depth + 1
                        j = j + 2
                    elseif text:byte(j) == 93 and text:byte(j+1) == 93 then
                        depth = depth - 1
                        j = j + 2
                    else
                        j = j + 1
                    end
                end
                i = j  -- 跳到 ]] 之后或文本末尾
            else
                -- 普通内部链接,保留
                table.insert(out, "[[")
                i = i + 2
            end
        else
            table.insert(out, text:sub(i,i))
            i = i + 1
        end
        
        if i > 1200 then break end  -- 极端安全阀
    end
    
    return table.concat(out)
end

-- ────────────────────────────────────────────────
-- 清理模板、参考文献、标题、列表等,得到纯文本摘要
-- ────────────────────────────────────────────────
local function cleanFinal(text)
    if not text or text == "" then return "" end
    
    text = text
        :gsub("{{[^{}]*}}", "")                  -- 模板(简单版,够用)
        :gsub("<ref.-</?ref[^>]*>", "")          -- 参考文献
        :gsub("<!%-%-.-%-%->", "")               -- HTML 注释
        :gsub("\n==+[^=]+==+", " ")              -- 各级标题
        :gsub("\n%s*[*#:;]+", " ")               -- 列表、定义列表
        :gsub("%s*\n%s*", " ")                   -- 所有换行 → 单个空格
        :gsub("%s+", " ")                        -- 压缩连续空格
    
    return mw.text.trim(text)
end

-- ────────────────────────────────────────────────
-- 主函数:生成带首图的简短摘要 + 标题
-- ────────────────────────────────────────────────
function p.getSummaryAndImage(frame)
    local pageName = mw.text.trim(frame.args[1] or "")
    if pageName == "" then return "" end
    
    local title = mw.title.new(pageName)
    if not title or not title.exists then return "" end
    
    local content = title:getContent()
    if not content or content == "" then return "" end
    
    -- 安全取前 300 个 unicode 字符(避免字节截断中文)
    local rawExcerpt = mw.ustring.sub(content, 1, 300)
    
    -- 提取第一张图片(优先标准写法)
    local firstImage = mw.ustring.match(rawExcerpt, "%[%[(File|文件|Image|图像):([^|%]]+)")
    
    -- 备选:更宽松匹配,但只接受文件类
    if not firstImage then
        local candidate = mw.ustring.match(rawExcerpt, "%[%[([^|%]]+)%]")
        if candidate then
            local ns = candidate:match("^([^:]+):")
            if ns and (ns:lower() == "file" or ns == "文件" or ns:lower() == "image" or ns == "图像") then
                firstImage = candidate
            end
        end
    end
    
    local step1     = stripImages(rawExcerpt) or ""
    local cleanText = cleanFinal(step1) or ""
    
    local targetLen = 120
    local summary   = cleanText:sub(1, targetLen * 2) or ""
    
    local ulen = mw.ustring.len(summary) or 0
    if ulen > targetLen then
        summary = mw.ustring.sub(summary, 1, targetLen) or ""
        
        -- 尽量避免截在中文单词/句子中间
        local rev = summary:reverse()
        local lastBreak = rev:find("[ %s%p,。!?;:,…]") or 1
        if lastBreak > 1 and lastBreak < 20 then
            summary = mw.ustring.sub(summary, 1, #summary - lastBreak + 1) or ""
        end
        
        summary = mw.text.trim(summary) .. "..."
    end
    
    summary = mw.text.trim(summary)
    
    -- ── HTML 输出 ───────────────────────────────────
    local container = mw.html.create('div')
        :css('margin-bottom', '25px')
        :css('display', 'flow-root')
    
    -- 标题
    container:tag('div')
        :css('font-size', '1.3em')
        :css('font-weight', 'bold')
        :css('margin-bottom', '6px')
        :css('border-bottom', '1px solid #eee')
        :wikitext('[[' .. pageName .. ']]')
    
    -- 内容区(右浮图 + 文字)
    local textDiv = container:tag('div')
        :css('line-height', '1.6')
        :css('color', '#222')
        :css('font-size', '14px')
    
    if firstImage and firstImage ~= "" then
        firstImage = mw.text.trim(firstImage)
        local img = '[[' .. firstImage .. '|120px|right|link=' .. pageName .. ']]'
        textDiv:wikitext(img)
    end
    
    textDiv:wikitext(summary)
    
    return tostring(container)
end

return p