打开/关闭菜单
60
73
29
1.1K
武外梗百科
打开/关闭外观设置菜单
打开/关闭个人菜单
未登录
未登录用户的IP地址会在进行任意编辑后公开展示。

Module:PageSummary:修订间差异

武外梗百科 爱国好学自强图新的百科全书
无编辑摘要
无编辑摘要
第1行: 第1行:
local p = {}
local p = {}


-- 核心清理函数:仅使用 gsub,不使用 while,从根源杜绝死循环
-- 核心剥离逻辑:状态机识别,绝对不回溯,绝对不死循环
local function safeStrip(text)
local function stripImages(text)
     if not text then return "" end
     if not text then return "" end
      
      
     -- 1. 移除引用和注释
     local out = {}
     text = string.gsub(text, "<ref.-</ref>", "")
     local i = 1
     text = string.gsub(text, "<ref.->", "")
     local len = mw.ustring.len(text)
    text = string.gsub(text, "<!%-%-.-%-%->", "")
      
      
     -- 2. 移除图片标签(分两步粉碎,处理掉带有一层链接的图片描述)
     while i <= len do
    local prefixes = {"[Ff]ile:", "[Ii]mage:", "文件:", "图像:"}
        local foundImage = false
    for _, prefix in ipairs(prefixes) do
        -- 检查是否是图片起始
        -- 第一步:爆破嵌套链接的图片 [[File:xxx|[[yyy]]]]
        local chunk = mw.ustring.sub(text, i, i + 10)
        text = string.gsub(text, "%[%[" .. prefix .. "[^%[%]]+ %[%[[^%[%]]+%]%] [^%[%]]+ %]%]", "")
        if chunk:match("^%[%[[Ff]ile:") or chunk:match("^%[%[[Ii]mage:") or
        -- 第二步:处理普通图片
          chunk:match("^%[%[文件:") or chunk:match("^%[%[图像:") then
         text = string.gsub(text, "%[%[" .. prefix .. ".-%]%]", "")
           
            -- 进入图片跳过模式
            local depth = 0
            local j = i
            while j <= len do
                local char2 = mw.ustring.sub(text, j, j + 1)
                if char2 == "[[" then
                    depth = depth + 1
                    j = j + 2
                elseif char2 == "]]" then
                    depth = depth - 1
                    j = j + 2
                    if depth <= 0 then
                        i = j -- 跳过整段图片
                        foundImage = true
                        break
                    end
                else
                    j = j + 1
                end
            end
        end
       
         if not foundImage then
            table.insert(out, mw.ustring.sub(text, i, i))
            i = i + 1
        end
       
        -- 安全熔断:如果处理太长,强制中断
        if i > 600 then break end
     end
     end
   
    return table.concat(out)
end


     -- 3. 移除模板 (不递归,只跑 3 次,保证性能)
local function cleanFinal(text)
     -- 移除模板
     text = string.gsub(text, "{{[^{}]+}}", "")
     text = string.gsub(text, "{{[^{}]+}}", "")
     text = string.gsub(text, "{{[^{}]+}}", "")
     text = string.gsub(text, "{{[^{}]+}}", "")
     text = string.gsub(text, "{{[^{}]+}}", "")
    -- 移除引用和注释
 
     text = string.gsub(text, "<ref.-</ref>", "")
     -- 4. 移除 Wiki 格式符号
    text = string.gsub(text, "<ref.->", "")
     text = string.gsub(text, "<!%-%-.-%-%->", "")
    -- 移除 Wiki 格式
     text = string.gsub(text, "\n==+.-==+", " ")
     text = string.gsub(text, "\n==+.-==+", " ")
     text = string.gsub(text, "\n%s*[*#:]+", " ")
     text = string.gsub(text, "\n%s*[*#:]+", " ")
    text = string.gsub(text, "&[Nn][Bb][Ss][Pp];", " ")
     text = string.gsub(text, "\n", " ")
     text = string.gsub(text, "\n", " ")
     text = string.gsub(text, "%s+", " ")
     text = string.gsub(text, "%s+", " ")
   
     return mw.text.trim(text)
     return mw.text.trim(text)
end
end
第40行: 第72行:
      
      
     local content = title:getContent()
     local content = title:getContent()
    -- 减少预加载长度到 400,从源头降低压力
     local rawExcerpt = mw.ustring.sub(content, 1, 400)
     local rawExcerpt = mw.ustring.sub(content, 1, 400)


     -- 提取第一张图(不带参数,直接拿名字)
     -- 提取第一张图(用于显示)
     local firstImage = string.match(rawExcerpt, "%[%[([Ff]ile:[^|%]%s]+)") or  
     local firstImage = string.match(rawExcerpt, "%[%[([Ff]ile:[^|%]%s]+)") or  
                       string.match(rawExcerpt, "%[%[(文件:[^|%]%s]+)")
                       string.match(rawExcerpt, "%[%[(文件:[^|%]%s]+)")


     -- 执行安全清理
     -- 执行两步清理
     local cleanText = safeStrip(rawExcerpt)
    local step1 = stripImages(rawExcerpt)
     local cleanText = cleanFinal(step1)


     -- 截断 150 字
     -- 截断
     local targetLen = 150
     local targetLen = 140
     local summary = mw.ustring.sub(cleanText, 1, targetLen)
     local summary = mw.ustring.sub(cleanText, 1, targetLen)
      
      
     -- 简单的截断修复
     -- 清理末尾断头链接
    summary = string.gsub(summary, "%[+$", "") -- 移除末尾残留的 [
   
     local _, opens = string.gsub(summary, "%[%[", "")
     local _, opens = string.gsub(summary, "%[%[", "")
     local _, closes = string.gsub(summary, "%]%]", "")
     local _, closes = string.gsub(summary, "%]%]", "")
     if opens > closes then
     if opens > closes then
        -- 链接不完整时,直接反向切掉最后一段,不进行向后搜索(向后搜索最耗资源)
         local lastOpen = string.find(summary:reverse(), "%[%[")
         local lastOpen = string.find(summary:reverse(), "%[%[")
         if lastOpen then
         if lastOpen then
第66行: 第95行:
         end
         end
     end
     end
   
 
     summary = mw.text.trim(summary)
     summary = mw.text.trim(summary)
     if mw.ustring.len(cleanText) > targetLen then
     if mw.ustring.len(cleanText) > targetLen then summary = summary .. "..." end
        summary = summary .. "..."
    end


     -- 渲染
     -- 渲染
     local container = mw.html.create('div'):css({['margin-bottom']='25px', ['display']='flow-root'})
     local container = mw.html.create('div'):css({['margin-bottom']='25px', ['display']='flow-root'})
     if firstImage then
     if firstImage then
         container:wikitext('[[' .. firstImage .. '|150px|right|link=' .. pageName .. ']]')
         container:wikitext('[[' .. firstImage .. '|130px|right|link=' .. pageName .. ']]')
     end
     end
     container:tag('div')
     container:tag('div')
         :css({['font-size']='1.35em', ['font-weight']='bold', ['margin-bottom']='6px'})
         :css({['font-size']='1.3em', ['font-weight']='bold', ['margin-bottom']='6px'})
         :wikitext('[[' .. pageName .. ']]')
         :wikitext('[[' .. pageName .. ']]')
     container:tag('div')
     container:tag('div')
         :css({['line-height']='1.6', ['color']='#202122'})
         :css({['line-height']='1.5', ['color']='#222'})
         :wikitext(summary)
         :wikitext(summary)



2025年12月21日 (日) 18:22的版本

此模块的文档可以在Module:PageSummary/doc创建

local p = {}

-- 核心剥离逻辑:状态机识别,绝对不回溯,绝对不死循环
local function stripImages(text)
    if not text then return "" end
    
    local out = {}
    local i = 1
    local len = mw.ustring.len(text)
    
    while i <= len do
        local foundImage = false
        -- 检查是否是图片起始
        local chunk = mw.ustring.sub(text, i, i + 10)
        if chunk:match("^%[%[[Ff]ile:") or chunk:match("^%[%[[Ii]mage:") or 
           chunk:match("^%[%[文件:") or chunk:match("^%[%[图像:") then
            
            -- 进入图片跳过模式
            local depth = 0
            local j = i
            while j <= len do
                local char2 = mw.ustring.sub(text, j, j + 1)
                if char2 == "[[" then
                    depth = depth + 1
                    j = j + 2
                elseif char2 == "]]" then
                    depth = depth - 1
                    j = j + 2
                    if depth <= 0 then
                        i = j -- 跳过整段图片
                        foundImage = true
                        break
                    end
                else
                    j = j + 1
                end
            end
        end
        
        if not foundImage then
            table.insert(out, mw.ustring.sub(text, i, i))
            i = i + 1
        end
        
        -- 安全熔断:如果处理太长,强制中断
        if i > 600 then break end
    end
    
    return table.concat(out)
end

local function cleanFinal(text)
    -- 移除模板
    text = string.gsub(text, "{{[^{}]+}}", "")
    text = string.gsub(text, "{{[^{}]+}}", "")
    -- 移除引用和注释
    text = string.gsub(text, "<ref.-</ref>", "")
    text = string.gsub(text, "<ref.->", "")
    text = string.gsub(text, "<!%-%-.-%-%->", "")
    -- 移除 Wiki 格式
    text = string.gsub(text, "\n==+.-==+", " ")
    text = string.gsub(text, "\n%s*[*#:]+", " ")
    text = string.gsub(text, "\n", " ")
    text = string.gsub(text, "%s+", " ")
    return mw.text.trim(text)
end

function p.getSummaryAndImage(frame)
    local pageName = frame.args[1] or ""
    local title = mw.title.new(pageName)
    if not title or not title.exists then return "" end
    
    local content = title:getContent()
    local rawExcerpt = mw.ustring.sub(content, 1, 400)

    -- 提取第一张图(用于显示)
    local firstImage = string.match(rawExcerpt, "%[%[([Ff]ile:[^|%]%s]+)") or 
                       string.match(rawExcerpt, "%[%[(文件:[^|%]%s]+)")

    -- 执行两步清理
    local step1 = stripImages(rawExcerpt)
    local cleanText = cleanFinal(step1)

    -- 截断
    local targetLen = 140
    local summary = mw.ustring.sub(cleanText, 1, targetLen)
    
    -- 清理末尾断头链接
    local _, opens = string.gsub(summary, "%[%[", "")
    local _, closes = string.gsub(summary, "%]%]", "")
    if opens > closes then
        local lastOpen = string.find(summary:reverse(), "%[%[")
        if lastOpen then
            summary = string.sub(summary, 1, #summary - lastOpen - 1)
        end
    end

    summary = mw.text.trim(summary)
    if mw.ustring.len(cleanText) > targetLen then summary = summary .. "..." end

    -- 渲染
    local container = mw.html.create('div'):css({['margin-bottom']='25px', ['display']='flow-root'})
    if firstImage then
        container:wikitext('[[' .. firstImage .. '|130px|right|link=' .. pageName .. ']]')
    end
    container:tag('div')
        :css({['font-size']='1.3em', ['font-weight']='bold', ['margin-bottom']='6px'})
        :wikitext('[[' .. pageName .. ']]')
    container:tag('div')
        :css({['line-height']='1.5', ['color']='#222'})
        :wikitext(summary)

    return tostring(container)
end

return p