↑ Up

Vampire---5.0.1.THM-Ref.s

View TPTP
Problem
Process solution in
SystemOnTSTP
Download .tgz
%------------------------------------------------------------------------------
% File     : Vampire---5.0.1
% Problem  : COM244_1 : TPTP v9.3.1. Released v9.3.0.
% Transfm  : none
% Format   : tptp:raw
% Command  : run_vampire /export/starexec/sandbox/benchmark/theBenchmark.p 300 THM

% Computer : n001.cluster.edu
% Model    : x86_64 x86_64
% CPU      : Intel(R) Xeon(R) CPU E5-2620 v4 2.10GHz
% Memory   : 8046.5625MB
% OS       : Linux 6.8.0-71-generic
% CPULimit : 300s
% WCLimit  : 300s
% DateTime : Tue Sep 29 09:39:52 AM UTC 2026

% Result   : Theorem 2.74s 1.16s
% Output   : Refutation 2.74s
% Verified : 
% SZS Type : Refutation
%            Derivation depth      :    7
%            Number of leaves      :    3
% Syntax   : Number of formulae    :   17 (   5 unt;   0 typ;   0 def)
%            Number of atoms       :   44 (  15 equ)
%            Maximal formula atoms :    5 (   2 avg)
%            Number of connectives :   50 (  23   ~;   7   |;  17   &)
%                                         (   0 <=>;   3  =>;   0  <=;   0 <~>)
%            Maximal formula depth :    8 (   4 avg)
%            Maximal term depth    :    2 (   1 avg)
%            Number of types       :    4 (   3 usr)
%            Number of type conns  :    0 (   0   >;   0   *;   0   +;   0  <<)
%            Number of predicates  :    7 (   5 usr;   1 prp; 0-3 aty)
%            Number of functors    :   52 (  52 usr;   7 con; 0-3 aty)
%            Number of variables   :   13 (   9   !;   4   ?;  13   :)

% Comments : 
%------------------------------------------------------------------------------
tff(type_def_5,type,
    vTerm: $tType ).

tff(type_def_6,type,
    vOptTerm: $tType ).

tff(type_def_7,type,
    vTy: $tType ).

tff(func_def_0,type,
    vNat: vTy ).

tff(func_def_1,type,
    vSucc: vTerm > vTerm ).

tff(func_def_2,type,
    vPred: vTerm > vTerm ).

tff(func_def_3,type,
    vsomeTerm: vTerm > vOptTerm ).

tff(func_def_4,type,
    vZero: vTerm ).

tff(func_def_5,type,
    vIszero: vTerm > vTerm ).

tff(func_def_6,type,
    vTrue: vTerm ).

tff(func_def_7,type,
    vFalse: vTerm ).

tff(func_def_8,type,
    vnoTerm: vOptTerm ).

tff(func_def_9,type,
    vPlus: ( vTerm * vTerm ) > vTerm ).

tff(func_def_10,type,
    vB: vTy ).

tff(func_def_11,type,
    vIfelse: ( vTerm * vTerm * vTerm ) > vTerm ).

tff(func_def_12,type,
    vreduce: vTerm > vOptTerm ).

tff(func_def_13,type,
    vplusop: ( vTerm * vTerm ) > vTerm ).

tff(func_def_14,type,
    vgetTerm: vOptTerm > vTerm ).

tff(func_def_15,type,
    sK1: vTy ).

tff(func_def_16,type,
    sK2: vTerm > vTerm ).

tff(func_def_17,type,
    sK3: vTerm > vTerm ).

tff(func_def_18,type,
    sK4: vTerm > vTerm ).

tff(func_def_19,type,
    sK5: vTerm > vTerm ).

tff(func_def_20,type,
    sK6: vTerm > vTerm ).

tff(func_def_21,type,
    sK7: vTerm > vTerm ).

tff(func_def_22,type,
    sK8: vTerm > vTerm ).

tff(func_def_23,type,
    sK9: vTerm > vTerm ).

tff(func_def_24,type,
    sK10: vTerm > vTerm ).

tff(func_def_25,type,
    sK11: vTerm > vTerm ).

tff(func_def_26,type,
    sK12: vTerm > vTerm ).

tff(func_def_27,type,
    sK13: vTerm > vTerm ).

tff(func_def_28,type,
    sK14: vTerm > vTerm ).

tff(func_def_29,type,
    sK15: vTerm > vTerm ).

tff(func_def_30,type,
    sK16: vTerm > vTerm ).

tff(func_def_31,type,
    sK17: vTerm > vTerm ).

tff(func_def_32,type,
    sK18: vTerm > vTerm ).

tff(func_def_33,type,
    sK19: vTerm > vTerm ).

tff(func_def_34,type,
    sK20: ( vTerm * vTerm * vTerm ) > vTerm ).

tff(func_def_35,type,
    sK21: ( vTerm * vTerm * vTerm ) > vTerm ).

tff(func_def_36,type,
    sK22: ( vTerm * vTerm * vTerm ) > vTerm ).

tff(func_def_37,type,
    sK23: ( vTerm * vTerm * vTerm ) > vTerm ).

tff(func_def_38,type,
    sK24: vOptTerm > vTerm ).

tff(func_def_39,type,
    sK25: vTerm > vTerm ).

tff(func_def_40,type,
    sK26: vTerm > vTerm ).

tff(func_def_41,type,
    sK27: vTerm > vTerm ).

tff(func_def_42,type,
    sK28: vTerm > vTerm ).

tff(func_def_43,type,
    sK29: vTerm > vTerm ).

tff(func_def_44,type,
    sK30: vTerm > vTerm ).

tff(func_def_45,type,
    sK31: vTerm > vTerm ).

tff(func_def_46,type,
    sK32: vTerm > vTerm ).

tff(func_def_47,type,
    sK33: vTerm > vTerm ).

tff(func_def_48,type,
    sK34: vTerm > vTerm ).

tff(func_def_49,type,
    sK35: vTerm > vTerm ).

tff(func_def_50,type,
    sK36: vTerm > vTerm ).

tff(func_def_51,type,
    sK37: vOptTerm > vTerm ).

tff(pred_def_1,type,
    visValue: vTerm > $o ).

tff(pred_def_2,type,
    visNV: vTerm > $o ).

tff(pred_def_3,type,
    visSomeTerm: vOptTerm > $o ).

tff(pred_def_4,type,
    vptchecksimple: ( vTerm * vTy ) > $o ).

tff(pred_def_5,type,
    sP0: ( vTerm * vTerm * vTerm ) > $o ).

tff(f41,axiom,
    visNV(vZero),
    file('/export/starexec/sandbox/benchmark/theBenchmark.p','isNV-0') ).

tff(f50,axiom,
    ! [X0: vTerm] :
      ( ~ visValue(X0)
     => ? [X1: vTerm] :
          ( ( X1 != vTrue )
          & ( X1 != vFalse )
          & ( X0 = X1 )
          & ~ visNV(X1) ) ),
    file('/export/starexec/sandbox/benchmark/theBenchmark.p','isValue-false-INV') ).

tff(f105,conjecture,
    ! [X0: vTy] :
      ( ( vptchecksimple(vZero,X0)
        & ~ visValue(vZero) )
     => ( vreduce(vZero) != vnoTerm ) ),
    file('/export/starexec/sandbox/benchmark/theBenchmark.p','Progress-Zero') ).

tff(f106,negated_conjecture,
    ~ ! [X0: vTy] :
        ( ( vptchecksimple(vZero,X0)
          & ~ visValue(vZero) )
       => ( vreduce(vZero) != vnoTerm ) ),
    inference(negated_conjecture,[status(cth)],[f105]) ).

tff(f109,plain,
    ? [X0: vTy] :
      ( ( vnoTerm = vreduce(vZero) )
      & vptchecksimple(vZero,X0)
      & ~ visValue(vZero) ),
    inference(ennf_transformation,[],[f106]) ).

tff(f110,plain,
    ? [X0: vTy] :
      ( ( vnoTerm = vreduce(vZero) )
      & vptchecksimple(vZero,X0)
      & ~ visValue(vZero) ),
    inference(flattening,[],[f109]) ).

tff(f111,plain,
    ! [X0: vTerm] :
      ( ? [X1: vTerm] :
          ( ( X1 != vTrue )
          & ( X1 != vFalse )
          & ( X0 = X1 )
          & ~ visNV(X1) )
      | visValue(X0) ),
    inference(ennf_transformation,[],[f50]) ).

tff(f153,plain,
    ( ( vnoTerm = vreduce(vZero) )
    & vptchecksimple(vZero,sK1)
    & ~ visValue(vZero) ),
    inference(skolemize,[status(esa),new_symbols(skolem,[sK1]),skolemize(X0,sK1)],[f110]) ).

tff(f154,plain,
    ! [X0: vTerm] :
      ( ( ( vTrue != sK2(X0) )
        & ( vFalse != sK2(X0) )
        & ( sK2(X0) = X0 )
        & ~ visNV(sK2(X0)) )
      | visValue(X0) ),
    inference(skolemize,[status(esa),new_symbols(skolem,[sK2]),skolemize(X1,sK2(X0))],[f111]) ).

tff(f171,plain,
    ~ visValue(vZero),
    inference(cnf_transformation,[],[f153]) ).

tff(f174,plain,
    ! [X0: vTerm] :
      ( ~ visNV(sK2(X0))
      | visValue(X0) ),
    inference(cnf_transformation,[],[f154]) ).

tff(f175,plain,
    ! [X0: vTerm] :
      ( ( sK2(X0) = X0 )
      | visValue(X0) ),
    inference(cnf_transformation,[],[f154]) ).

tff(f217,plain,
    visNV(vZero),
    inference(cnf_transformation,[],[f41]) ).

tff(f267,plain,
    ! [X0: vTerm] :
      ( ~ visNV(X0)
      | visValue(X0)
      | visValue(X0) ),
    inference(superposition,[],[f174,f175]) ).

tff(f268,plain,
    ! [X0: vTerm] :
      ( visValue(X0)
      | ~ visNV(X0) ),
    inference(duplicate_literal_removal,[],[f267]) ).

tff(f269,plain,
    ~ visNV(vZero),
    inference(resolution,[],[f268,f171]) ).

tff(f270,plain,
    $false,
    inference(forward_subsumption_resolution,[],[f269,f217]) ).

%------------------------------------------------------------------------------
%----ORIGINAL SYSTEM OUTPUT
% 0.00/0.03  % Problem  : COM244_1 : TPTP v9.3.1. Released v9.3.0.
% 0.00/0.05  % Command  : run_vampire /export/starexec/sandbox/benchmark/theBenchmark.p 300 THM
% 0.08/0.17  % Computer : n001.cluster.edu
% 0.08/0.17  % Model    : x86_64 x86_64
% 0.08/0.17  % CPU      : Intel(R) Xeon(R) CPU E5-2620 v4 @ 2.10GHz
% 0.08/0.17  % Memory   : 8046.5625MB
% 0.08/0.17  % OS       : Linux 6.8.0-71-generic
% 0.08/0.17  % CPULimit : 300
% 0.08/0.17  % WCLimit  : 300
% 0.08/0.17  % DateTime : Mon Sep 28 22:08:19 UTC 2026
% 0.08/0.17  % CPUTime  : 
% 0.08/0.17  Running run_vampire /export/starexec/sandbox/benchmark/theBenchmark.p 300 THM
% 0.08/0.21  Running first-order theorem proving
% 0.08/0.21  Running: /export/starexec/sandbox/solver/bin/vampire --input_syntax tptp --output_axiom_names on --mode casc -m 16384 --cores 7 -t 300 /export/starexec/sandbox/benchmark/theBenchmark.p
% 2.74/1.16  % (797704)Detected formulas, will run a generic FOF schedule.
% 2.74/1.16  % (797814)dis-1010_2:3_sil=16000:sp=reverse_frequency:random_seed=3145339436:i=119:av=off:ss=axioms_2999 on theBenchmark for (2999ds/119Mi)
% 2.74/1.16  % (797814)First to succeed.
% 2.74/1.16  % (797814)Solution written to "/export/starexec/sandbox/tmp/vampire-proof-797704"
% 2.74/1.16  % (797813)lrs+1010_1_to=lpo:sil=32000:sos=on:spb=goal_then_units:bce=on:random_seed=2013023070:i=109:sd=1:ins=1:gsp=on:ss=axioms_2999 on theBenchmark for (2999ds/109Mi)
% 2.74/1.16  % (797810)lrs+10_1_ncem=casc2026/models/loop8.pt:sil=128000:tgt=full:npcc=on:drc=off:sp=weighted_frequency:spb=goal:fd=preordered:foolp=on:random_seed=3733875117:i=141193_2999 on theBenchmark for (2999ds/141193Mi)
% 2.74/1.16  % (797812)lrs+1010_1_anc=all:sfv=off:to=kbo:ncem=casc2026/models/loop7.pt:sil=128000:npcc=on:prc=on:sos=all:bsr=unit_only:sac=on:random_seed=761956865:i=141695:sd=1:nm=32:gsp=on:ss=included_2999 on theBenchmark for (2999ds/141695Mi)
% 2.74/1.16  % (797811)lrs+11_1_ncem=casc2026/models/loop8.pt:sil=128000:npcc=on:lma=off:spb=units:urr=ec_only:bce=on:s2agt=64:updr=off:random_seed=52887203:i=134677:sd=20:aac=none:nm=16:ss=included:sgt=10_2999 on theBenchmark for (2999ds/134677Mi)
% 2.74/1.16  % (797813)Refutation not found, incomplete strategy
% 2.74/1.16  % (797813)------------------------------
% 2.74/1.16  % (797813)Version: Vampire 5.0.1 (Release build, commit ea8961452 on 2026-07-16 15:14:34 +0200)
% 2.74/1.16  % (797813)Linked with Z3 4.14.0.0 3c47fd96cf5645d0c42b2c819d9e9a84380aa721 z3-4.8.4-9178-g3c47fd96c
% 2.74/1.16  % (797813)CaDiCaL version: 2.1.3
% 2.74/1.16  % (797813)Termination reason: Refutation not found, incomplete strategy
% 2.74/1.16  % (797813)Time elapsed: 0.003 s
% 2.74/1.16  % (797813)Peak memory usage: 88 MB
% 2.74/1.16  % (797813)Instructions burned: 3 (million)
% 2.74/1.16  % (797816)dis-21_1_sil=8000:lcm=predicate:random_seed=1502320686:st=5:avsq=on:i=129:avsqr=1,16:sd=3:aac=none:ep=RS:fsr=off:ss=included_2999 on theBenchmark for (2999ds/129Mi)
% 2.74/1.16  % (797815)dis-1011_1_sil=16000:fde=unused:s2agt=70:random_seed=2894824516:s2a=on:i=139:gtg=position_2999 on theBenchmark for (2999ds/139Mi)
% 2.74/1.16  % (797815)Also succeeded, but the first one will report.
% 2.74/1.16  % (797816)Also succeeded, but the first one will report.
% 2.74/1.16  % (797814)Refutation found. Thanks to Tanya!
% 2.74/1.16  % SZS status Theorem for theBenchmark
% 2.74/1.16  % SZS output start Proof for theBenchmark
% See solution above
% 2.74/1.16  % (797814)------------------------------
% 2.74/1.16  % (797814)Version: Vampire 5.0.1 (Release build, commit ea8961452 on 2026-07-16 15:14:34 +0200)
% 2.74/1.16  % (797814)Linked with Z3 4.14.0.0 3c47fd96cf5645d0c42b2c819d9e9a84380aa721 z3-4.8.4-9178-g3c47fd96c
% 2.74/1.16  % (797814)CaDiCaL version: 2.1.3
% 2.74/1.16  % (797814)Termination reason: Refutation
% 2.74/1.16  % (797814)Time elapsed: 0.003 s
% 2.74/1.16  % (797814)Peak memory usage: 88 MB
% 2.74/1.16  % (797814)Instructions burned: 4 (million)
% 2.74/1.16  % (797814)------------------------------
% 2.74/1.16  % (797814)------------------------------
% 2.74/1.16  % (797704)Success in time 0.301 s
% 2.74/1.16  % Vampire exiting
%------------------------------------------------------------------------------